| { | |
| "d_model": 128, | |
| "nhead": 4, | |
| "num_layers": 4, | |
| "vocab_size": 257, | |
| "avg_loss": 0.05360993637525373, | |
| "avg_bpb": 0.07734278935095133 | |
| } |
| { | |
| "d_model": 128, | |
| "nhead": 4, | |
| "num_layers": 4, | |
| "vocab_size": 257, | |
| "avg_loss": 0.05360993637525373, | |
| "avg_bpb": 0.07734278935095133 | |
| } |