{ "best_metric": null, "best_model_checkpoint": null, "epoch": 3.0, "eval_steps": 500, "global_step": 243, "is_hyper_param_search": false, "is_local_process_zero": true, "is_world_process_zero": true, "log_history": [ { "epoch": 1.0, "eval_B-Claim": { "f1-score": 0.37160751565762, "precision": 0.42788461538461536, "recall": 0.3284132841328413, "support": 271.0 }, "eval_B-MajorClaim": { "f1-score": 0.3578947368421052, "precision": 0.6666666666666666, "recall": 0.2446043165467626, "support": 139.0 }, "eval_B-Premise": { "f1-score": 0.8640915593705293, "precision": 0.7895424836601307, "recall": 0.9541864139020537, "support": 633.0 }, "eval_I-Claim": { "f1-score": 0.5003402749421533, "precision": 0.5493126120741183, "recall": 0.4593851537115721, "support": 4001.0 }, "eval_I-MajorClaim": { "f1-score": 0.7718093699515347, "precision": 0.6502211636611093, "recall": 0.9493293591654247, "support": 2013.0 }, "eval_I-Premise": { "f1-score": 0.875016720916752, "precision": 0.8846812731043188, "recall": 0.865561044460127, "support": 11336.0 }, "eval_O": { "f1-score": 0.9992483530087988, "precision": 0.9995577178239717, "recall": 0.998939179632249, "support": 11312.0 }, "eval_accuracy": 0.8614038040733883, "eval_loss": 0.31713685393333435, "eval_macro avg": { "f1-score": 0.6771440758127848, "precision": 0.7096952189107044, "recall": 0.685774107364433, "support": 29705.0 }, "eval_runtime": 4.8338, "eval_samples_per_second": 16.55, "eval_steps_per_second": 2.069, "eval_weighted avg": { "f1-score": 0.8576207231627551, "precision": 0.8601529227027923, "recall": 0.8614038040733883, "support": 29705.0 }, "step": 81 }, { "epoch": 2.0, "eval_B-Claim": { "f1-score": 0.4708624708624709, "precision": 0.6392405063291139, "recall": 0.3726937269372694, "support": 271.0 }, "eval_B-MajorClaim": { "f1-score": 0.796875, "precision": 0.8717948717948718, "recall": 0.7338129496402878, "support": 139.0 }, "eval_B-Premise": { "f1-score": 0.8736616702355461, "precision": 0.796875, "recall": 0.966824644549763, "support": 633.0 }, "eval_I-Claim": { "f1-score": 0.5100589925881107, "precision": 0.6459770114942529, "recall": 0.4213946513371657, "support": 4001.0 }, "eval_I-MajorClaim": { "f1-score": 0.8401387776888176, "precision": 0.9077277970011534, "recall": 0.7819175360158966, "support": 2013.0 }, "eval_I-Premise": { "f1-score": 0.8912891699864469, "precision": 0.8338584492430646, "recall": 0.9572159491884262, "support": 11336.0 }, "eval_O": { "f1-score": 0.9996904982977407, "precision": 1.0, "recall": 0.9993811881188119, "support": 11312.0 }, "eval_accuracy": 0.8830499915839084, "eval_loss": 0.2966194748878479, "eval_macro avg": { "f1-score": 0.7689395113798762, "precision": 0.8136390908374939, "recall": 0.7476058065410885, "support": 29705.0 }, "eval_runtime": 4.8625, "eval_samples_per_second": 16.452, "eval_steps_per_second": 2.057, "eval_weighted avg": { "f1-score": 0.8731020208182413, "precision": 0.874440834821272, "recall": 0.8830499915839084, "support": 29705.0 }, "step": 162 }, { "epoch": 3.0, "eval_B-Claim": { "f1-score": 0.6085192697768763, "precision": 0.6756756756756757, "recall": 0.5535055350553506, "support": 271.0 }, "eval_B-MajorClaim": { "f1-score": 0.8571428571428571, "precision": 0.851063829787234, "recall": 0.8633093525179856, "support": 139.0 }, "eval_B-Premise": { "f1-score": 0.8834729626808834, "precision": 0.8529411764705882, "recall": 0.9162717219589257, "support": 633.0 }, "eval_I-Claim": { "f1-score": 0.5764474423833614, "precision": 0.6584269662921348, "recall": 0.5126218445388653, "support": 4001.0 }, "eval_I-MajorClaim": { "f1-score": 0.8581151832460733, "precision": 0.9070282235749861, "recall": 0.8142076502732241, "support": 2013.0 }, "eval_I-Premise": { "f1-score": 0.8959744247675935, "precision": 0.8563158317922328, "recall": 0.939484827099506, "support": 11336.0 }, "eval_O": { "f1-score": 0.9996020340481981, "precision": 1.0, "recall": 0.9992043847241867, "support": 11312.0 }, "eval_accuracy": 0.8918700555462044, "eval_loss": 0.2552729547023773, "eval_macro avg": { "f1-score": 0.811324882006549, "precision": 0.8287788147989789, "recall": 0.7998007594525776, "support": 29705.0 }, "eval_runtime": 4.8422, "eval_samples_per_second": 16.522, "eval_steps_per_second": 2.065, "eval_weighted avg": { "f1-score": 0.8867633844066056, "precision": 0.886070631898416, "recall": 0.8918700555462044, "support": 29705.0 }, "step": 243 } ], "logging_steps": 500, "max_steps": 4050, "num_input_tokens_seen": 0, "num_train_epochs": 50, "save_steps": 500, "total_flos": 431372438154000.0, "train_batch_size": 4, "trial_name": null, "trial_params": null }