GPT2-NECK-SWEEP_20240119-032414-5b68c_dataset_name-rand_hidden_idxs-11_hidden_lb-0_neck_cls-lstm_pretrained-1_token_lb-0_epoch=02-val_self_loss=8.28.ckpt