GPT2-NECK-SWEEP_20240120-020159-bDb7f_dataset_name-msmarco_hidden_idxs-12_hidden_lb-0_neck_cls-lstm_pretrained-1_token_lb-0_epoch=02-val_self_loss=4.18.ckpt