| { | |
| "base_model": "meta-llama/Llama-3.2-3B-Instruct", | |
| "properties": [ | |
| "Egc", | |
| "Egb", | |
| "Eea", | |
| "Ei", | |
| "EPS", | |
| "Nc", | |
| "Xc", | |
| "Eat" | |
| ], | |
| "split_seed": 42, | |
| "train_examples": 152406, | |
| "assistant_only_loss": true, | |
| "method": "LoRA", | |
| "max_attachments": 2, | |
| "answer": "rules", | |
| "split_method": "component", | |
| "test_fraction": 0.2, | |
| "data_sha256": "622a31c09513369e2111ebb9919c58c5c2b5f2e75c878d4d27e336b339270405", | |
| "metrics": { | |
| "train_runtime": 17305.0779, | |
| "train_samples_per_second": 26.421, | |
| "train_steps_per_second": 0.413, | |
| "total_flos": 1.5334674210110833e+18, | |
| "train_loss": 0.028584727216574547, | |
| "epoch": 2.9993701448666807, | |
| "step": 7143 | |
| } | |
| } |