apwic commited on
Commit
016b28c
·
verified ·
1 Parent(s): 0b15301

End of training

Browse files
README.md CHANGED
@@ -1,4 +1,6 @@
1
  ---
 
 
2
  license: mit
3
  base_model: indolem/indobert-base-uncased
4
  tags:
 
1
  ---
2
+ language:
3
+ - id
4
  license: mit
5
  base_model: indolem/indobert-base-uncased
6
  tags:
all_results.json CHANGED
@@ -17,10 +17,10 @@
17
  "eval_overall_f1": 0.9390862944162437,
18
  "eval_overall_precision": 0.9343434343434344,
19
  "eval_overall_recall": 0.9438775510204082,
20
- "eval_runtime": 0.3078,
21
  "eval_samples": 170,
22
- "eval_samples_per_second": 552.251,
23
- "eval_steps_per_second": 9.746,
24
  "predict_LOCATION_f1": 0.8884615384615384,
25
  "predict_LOCATION_number": 261,
26
  "predict_LOCATION_precision": 0.8918918918918919,
@@ -38,12 +38,12 @@
38
  "predict_overall_f1": 0.9173789173789174,
39
  "predict_overall_precision": 0.9061913696060038,
40
  "predict_overall_recall": 0.9288461538461539,
41
- "predict_runtime": 0.7618,
42
- "predict_samples_per_second": 557.886,
43
- "predict_steps_per_second": 9.189,
44
  "train_loss": 0.03753832100580136,
45
- "train_runtime": 561.7168,
46
  "train_samples": 1531,
47
- "train_samples_per_second": 272.557,
48
- "train_steps_per_second": 17.09
49
  }
 
17
  "eval_overall_f1": 0.9390862944162437,
18
  "eval_overall_precision": 0.9343434343434344,
19
  "eval_overall_recall": 0.9438775510204082,
20
+ "eval_runtime": 0.3104,
21
  "eval_samples": 170,
22
+ "eval_samples_per_second": 547.675,
23
+ "eval_steps_per_second": 9.665,
24
  "predict_LOCATION_f1": 0.8884615384615384,
25
  "predict_LOCATION_number": 261,
26
  "predict_LOCATION_precision": 0.8918918918918919,
 
38
  "predict_overall_f1": 0.9173789173789174,
39
  "predict_overall_precision": 0.9061913696060038,
40
  "predict_overall_recall": 0.9288461538461539,
41
+ "predict_runtime": 0.7599,
42
+ "predict_samples_per_second": 559.268,
43
+ "predict_steps_per_second": 9.211,
44
  "train_loss": 0.03753832100580136,
45
+ "train_runtime": 578.6389,
46
  "train_samples": 1531,
47
+ "train_samples_per_second": 264.586,
48
+ "train_steps_per_second": 16.591
49
  }
eval_results.json CHANGED
@@ -17,8 +17,8 @@
17
  "eval_overall_f1": 0.9390862944162437,
18
  "eval_overall_precision": 0.9343434343434344,
19
  "eval_overall_recall": 0.9438775510204082,
20
- "eval_runtime": 0.3078,
21
  "eval_samples": 170,
22
- "eval_samples_per_second": 552.251,
23
- "eval_steps_per_second": 9.746
24
  }
 
17
  "eval_overall_f1": 0.9390862944162437,
18
  "eval_overall_precision": 0.9343434343434344,
19
  "eval_overall_recall": 0.9438775510204082,
20
+ "eval_runtime": 0.3104,
21
  "eval_samples": 170,
22
+ "eval_samples_per_second": 547.675,
23
+ "eval_steps_per_second": 9.665
24
  }
predict_results.json CHANGED
@@ -16,7 +16,7 @@
16
  "predict_overall_f1": 0.9173789173789174,
17
  "predict_overall_precision": 0.9061913696060038,
18
  "predict_overall_recall": 0.9288461538461539,
19
- "predict_runtime": 0.7618,
20
- "predict_samples_per_second": 557.886,
21
- "predict_steps_per_second": 9.189
22
  }
 
16
  "predict_overall_f1": 0.9173789173789174,
17
  "predict_overall_precision": 0.9061913696060038,
18
  "predict_overall_recall": 0.9288461538461539,
19
+ "predict_runtime": 0.7599,
20
+ "predict_samples_per_second": 559.268,
21
+ "predict_steps_per_second": 9.211
22
  }
runs/Jun04_11-10-22_a358b85c7679/events.out.tfevents.1717500015.a358b85c7679.641397.1 ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c032b964c6ff3324bd3c97367492e88e5b408f99c0f0917538889f5e1dd09342
3
+ size 1305
train_results.json CHANGED
@@ -1,8 +1,8 @@
1
  {
2
  "epoch": 100.0,
3
  "train_loss": 0.03753832100580136,
4
- "train_runtime": 561.7168,
5
  "train_samples": 1531,
6
- "train_samples_per_second": 272.557,
7
- "train_steps_per_second": 17.09
8
  }
 
1
  {
2
  "epoch": 100.0,
3
  "train_loss": 0.03753832100580136,
4
+ "train_runtime": 578.6389,
5
  "train_samples": 1531,
6
+ "train_samples_per_second": 264.586,
7
+ "train_steps_per_second": 16.591
8
  }
trainer_state.json CHANGED
@@ -34,9 +34,9 @@
34
  "eval_overall_f1": 0.23699421965317918,
35
  "eval_overall_precision": 0.2733333333333333,
36
  "eval_overall_recall": 0.20918367346938777,
37
- "eval_runtime": 0.3003,
38
- "eval_samples_per_second": 566.191,
39
- "eval_steps_per_second": 9.992,
40
  "step": 96
41
  },
42
  {
@@ -65,9 +65,9 @@
65
  "eval_overall_f1": 0.585305105853051,
66
  "eval_overall_precision": 0.5717761557177615,
67
  "eval_overall_recall": 0.5994897959183674,
68
- "eval_runtime": 0.3067,
69
- "eval_samples_per_second": 554.267,
70
- "eval_steps_per_second": 9.781,
71
  "step": 192
72
  },
73
  {
@@ -96,9 +96,9 @@
96
  "eval_overall_f1": 0.8114143920595533,
97
  "eval_overall_precision": 0.7898550724637681,
98
  "eval_overall_recall": 0.8341836734693877,
99
- "eval_runtime": 0.2971,
100
- "eval_samples_per_second": 572.18,
101
- "eval_steps_per_second": 10.097,
102
  "step": 288
103
  },
104
  {
@@ -127,9 +127,9 @@
127
  "eval_overall_f1": 0.8335388409371147,
128
  "eval_overall_precision": 0.8066825775656324,
129
  "eval_overall_recall": 0.8622448979591837,
130
- "eval_runtime": 0.2979,
131
- "eval_samples_per_second": 570.712,
132
- "eval_steps_per_second": 10.071,
133
  "step": 384
134
  },
135
  {
@@ -158,9 +158,9 @@
158
  "eval_overall_f1": 0.8955223880597014,
159
  "eval_overall_precision": 0.8737864077669902,
160
  "eval_overall_recall": 0.9183673469387755,
161
- "eval_runtime": 0.2919,
162
- "eval_samples_per_second": 582.404,
163
- "eval_steps_per_second": 10.278,
164
  "step": 480
165
  },
166
  {
@@ -189,9 +189,9 @@
189
  "eval_overall_f1": 0.8664987405541562,
190
  "eval_overall_precision": 0.8557213930348259,
191
  "eval_overall_recall": 0.8775510204081632,
192
- "eval_runtime": 0.2942,
193
- "eval_samples_per_second": 577.9,
194
- "eval_steps_per_second": 10.198,
195
  "step": 576
196
  },
197
  {
@@ -220,9 +220,9 @@
220
  "eval_overall_f1": 0.9052369077306733,
221
  "eval_overall_precision": 0.8853658536585366,
222
  "eval_overall_recall": 0.9260204081632653,
223
- "eval_runtime": 0.2999,
224
- "eval_samples_per_second": 566.894,
225
- "eval_steps_per_second": 10.004,
226
  "step": 672
227
  },
228
  {
@@ -251,9 +251,9 @@
251
  "eval_overall_f1": 0.9081761006289308,
252
  "eval_overall_precision": 0.8957816377171216,
253
  "eval_overall_recall": 0.9209183673469388,
254
- "eval_runtime": 0.2943,
255
- "eval_samples_per_second": 577.608,
256
- "eval_steps_per_second": 10.193,
257
  "step": 768
258
  },
259
  {
@@ -282,9 +282,9 @@
282
  "eval_overall_f1": 0.9051833122629582,
283
  "eval_overall_precision": 0.8972431077694235,
284
  "eval_overall_recall": 0.9132653061224489,
285
- "eval_runtime": 0.292,
286
- "eval_samples_per_second": 582.255,
287
- "eval_steps_per_second": 10.275,
288
  "step": 864
289
  },
290
  {
@@ -313,9 +313,9 @@
313
  "eval_overall_f1": 0.9120603015075376,
314
  "eval_overall_precision": 0.8985148514851485,
315
  "eval_overall_recall": 0.9260204081632653,
316
- "eval_runtime": 0.2942,
317
- "eval_samples_per_second": 577.914,
318
- "eval_steps_per_second": 10.198,
319
  "step": 960
320
  },
321
  {
@@ -344,9 +344,9 @@
344
  "eval_overall_f1": 0.9345088161209069,
345
  "eval_overall_precision": 0.9228855721393034,
346
  "eval_overall_recall": 0.9464285714285714,
347
- "eval_runtime": 0.2952,
348
- "eval_samples_per_second": 575.922,
349
- "eval_steps_per_second": 10.163,
350
  "step": 1056
351
  },
352
  {
@@ -375,9 +375,9 @@
375
  "eval_overall_f1": 0.9371859296482412,
376
  "eval_overall_precision": 0.9232673267326733,
377
  "eval_overall_recall": 0.951530612244898,
378
- "eval_runtime": 0.294,
379
- "eval_samples_per_second": 578.323,
380
- "eval_steps_per_second": 10.206,
381
  "step": 1152
382
  },
383
  {
@@ -406,9 +406,9 @@
406
  "eval_overall_f1": 0.9319899244332494,
407
  "eval_overall_precision": 0.9203980099502488,
408
  "eval_overall_recall": 0.9438775510204082,
409
- "eval_runtime": 0.2915,
410
- "eval_samples_per_second": 583.134,
411
- "eval_steps_per_second": 10.291,
412
  "step": 1248
413
  },
414
  {
@@ -437,9 +437,9 @@
437
  "eval_overall_f1": 0.9362244897959183,
438
  "eval_overall_precision": 0.9362244897959183,
439
  "eval_overall_recall": 0.9362244897959183,
440
- "eval_runtime": 0.2966,
441
- "eval_samples_per_second": 573.097,
442
- "eval_steps_per_second": 10.113,
443
  "step": 1344
444
  },
445
  {
@@ -468,9 +468,9 @@
468
  "eval_overall_f1": 0.9287531806615775,
469
  "eval_overall_precision": 0.9263959390862944,
470
  "eval_overall_recall": 0.9311224489795918,
471
- "eval_runtime": 0.2962,
472
- "eval_samples_per_second": 573.872,
473
- "eval_steps_per_second": 10.127,
474
  "step": 1440
475
  },
476
  {
@@ -499,9 +499,9 @@
499
  "eval_overall_f1": 0.9457755359394704,
500
  "eval_overall_precision": 0.9351620947630923,
501
  "eval_overall_recall": 0.9566326530612245,
502
- "eval_runtime": 0.2955,
503
- "eval_samples_per_second": 575.309,
504
- "eval_steps_per_second": 10.153,
505
  "step": 1536
506
  },
507
  {
@@ -530,9 +530,9 @@
530
  "eval_overall_f1": 0.9445843828715365,
531
  "eval_overall_precision": 0.9328358208955224,
532
  "eval_overall_recall": 0.9566326530612245,
533
- "eval_runtime": 0.2934,
534
- "eval_samples_per_second": 579.387,
535
- "eval_steps_per_second": 10.224,
536
  "step": 1632
537
  },
538
  {
@@ -561,9 +561,9 @@
561
  "eval_overall_f1": 0.9489795918367347,
562
  "eval_overall_precision": 0.9489795918367347,
563
  "eval_overall_recall": 0.9489795918367347,
564
- "eval_runtime": 0.2953,
565
- "eval_samples_per_second": 575.721,
566
- "eval_steps_per_second": 10.16,
567
  "step": 1728
568
  },
569
  {
@@ -592,9 +592,9 @@
592
  "eval_overall_f1": 0.9516539440203563,
593
  "eval_overall_precision": 0.949238578680203,
594
  "eval_overall_recall": 0.9540816326530612,
595
- "eval_runtime": 0.2917,
596
- "eval_samples_per_second": 582.72,
597
- "eval_steps_per_second": 10.283,
598
  "step": 1824
599
  },
600
  {
@@ -623,9 +623,9 @@
623
  "eval_overall_f1": 0.9333333333333335,
624
  "eval_overall_precision": 0.9205955334987593,
625
  "eval_overall_recall": 0.9464285714285714,
626
- "eval_runtime": 0.2956,
627
- "eval_samples_per_second": 575.138,
628
- "eval_steps_per_second": 10.149,
629
  "step": 1920
630
  },
631
  {
@@ -654,9 +654,9 @@
654
  "eval_overall_f1": 0.9402795425667091,
655
  "eval_overall_precision": 0.9367088607594937,
656
  "eval_overall_recall": 0.9438775510204082,
657
- "eval_runtime": 0.2935,
658
- "eval_samples_per_second": 579.13,
659
- "eval_steps_per_second": 10.22,
660
  "step": 2016
661
  },
662
  {
@@ -685,9 +685,9 @@
685
  "eval_overall_f1": 0.9213483146067415,
686
  "eval_overall_precision": 0.902200488997555,
687
  "eval_overall_recall": 0.9413265306122449,
688
- "eval_runtime": 0.2947,
689
- "eval_samples_per_second": 576.887,
690
- "eval_steps_per_second": 10.18,
691
  "step": 2112
692
  },
693
  {
@@ -716,9 +716,9 @@
716
  "eval_overall_f1": 0.9444444444444445,
717
  "eval_overall_precision": 0.935,
718
  "eval_overall_recall": 0.9540816326530612,
719
- "eval_runtime": 0.2954,
720
- "eval_samples_per_second": 575.532,
721
- "eval_steps_per_second": 10.156,
722
  "step": 2208
723
  },
724
  {
@@ -747,9 +747,9 @@
747
  "eval_overall_f1": 0.929113924050633,
748
  "eval_overall_precision": 0.9221105527638191,
749
  "eval_overall_recall": 0.9362244897959183,
750
- "eval_runtime": 0.2927,
751
- "eval_samples_per_second": 580.837,
752
- "eval_steps_per_second": 10.25,
753
  "step": 2304
754
  },
755
  {
@@ -778,9 +778,9 @@
778
  "eval_overall_f1": 0.9438775510204082,
779
  "eval_overall_precision": 0.9438775510204082,
780
  "eval_overall_recall": 0.9438775510204082,
781
- "eval_runtime": 0.2936,
782
- "eval_samples_per_second": 578.97,
783
- "eval_steps_per_second": 10.217,
784
  "step": 2400
785
  },
786
  {
@@ -809,9 +809,9 @@
809
  "eval_overall_f1": 0.9390862944162437,
810
  "eval_overall_precision": 0.9343434343434344,
811
  "eval_overall_recall": 0.9438775510204082,
812
- "eval_runtime": 0.2929,
813
- "eval_samples_per_second": 580.497,
814
- "eval_steps_per_second": 10.244,
815
  "step": 2496
816
  },
817
  {
@@ -840,9 +840,9 @@
840
  "eval_overall_f1": 0.9385194479297364,
841
  "eval_overall_precision": 0.9234567901234568,
842
  "eval_overall_recall": 0.9540816326530612,
843
- "eval_runtime": 0.2929,
844
- "eval_samples_per_second": 580.435,
845
- "eval_steps_per_second": 10.243,
846
  "step": 2592
847
  },
848
  {
@@ -871,9 +871,9 @@
871
  "eval_overall_f1": 0.9368686868686869,
872
  "eval_overall_precision": 0.9275,
873
  "eval_overall_recall": 0.9464285714285714,
874
- "eval_runtime": 0.293,
875
- "eval_samples_per_second": 580.134,
876
- "eval_steps_per_second": 10.238,
877
  "step": 2688
878
  },
879
  {
@@ -902,9 +902,9 @@
902
  "eval_overall_f1": 0.9343434343434343,
903
  "eval_overall_precision": 0.925,
904
  "eval_overall_recall": 0.9438775510204082,
905
- "eval_runtime": 0.2935,
906
- "eval_samples_per_second": 579.251,
907
- "eval_steps_per_second": 10.222,
908
  "step": 2784
909
  },
910
  {
@@ -933,9 +933,9 @@
933
  "eval_overall_f1": 0.9433962264150944,
934
  "eval_overall_precision": 0.9305210918114144,
935
  "eval_overall_recall": 0.9566326530612245,
936
- "eval_runtime": 0.294,
937
- "eval_samples_per_second": 578.203,
938
- "eval_steps_per_second": 10.204,
939
  "step": 2880
940
  },
941
  {
@@ -964,9 +964,9 @@
964
  "eval_overall_f1": 0.9389312977099236,
965
  "eval_overall_precision": 0.9365482233502538,
966
  "eval_overall_recall": 0.9413265306122449,
967
- "eval_runtime": 0.2962,
968
- "eval_samples_per_second": 573.976,
969
- "eval_steps_per_second": 10.129,
970
  "step": 2976
971
  },
972
  {
@@ -995,9 +995,9 @@
995
  "eval_overall_f1": 0.9328263624841572,
996
  "eval_overall_precision": 0.9269521410579346,
997
  "eval_overall_recall": 0.9387755102040817,
998
- "eval_runtime": 0.2936,
999
- "eval_samples_per_second": 579.067,
1000
- "eval_steps_per_second": 10.219,
1001
  "step": 3072
1002
  },
1003
  {
@@ -1026,9 +1026,9 @@
1026
  "eval_overall_f1": 0.9491094147582698,
1027
  "eval_overall_precision": 0.9467005076142132,
1028
  "eval_overall_recall": 0.951530612244898,
1029
- "eval_runtime": 0.2959,
1030
- "eval_samples_per_second": 574.521,
1031
- "eval_steps_per_second": 10.139,
1032
  "step": 3168
1033
  },
1034
  {
@@ -1057,9 +1057,9 @@
1057
  "eval_overall_f1": 0.9444444444444445,
1058
  "eval_overall_precision": 0.935,
1059
  "eval_overall_recall": 0.9540816326530612,
1060
- "eval_runtime": 0.2934,
1061
- "eval_samples_per_second": 579.495,
1062
- "eval_steps_per_second": 10.226,
1063
  "step": 3264
1064
  },
1065
  {
@@ -1088,9 +1088,9 @@
1088
  "eval_overall_f1": 0.9429657794676806,
1089
  "eval_overall_precision": 0.9370277078085643,
1090
  "eval_overall_recall": 0.9489795918367347,
1091
- "eval_runtime": 0.2964,
1092
- "eval_samples_per_second": 573.583,
1093
- "eval_steps_per_second": 10.122,
1094
  "step": 3360
1095
  },
1096
  {
@@ -1119,9 +1119,9 @@
1119
  "eval_overall_f1": 0.9211195928753181,
1120
  "eval_overall_precision": 0.9187817258883249,
1121
  "eval_overall_recall": 0.923469387755102,
1122
- "eval_runtime": 0.2932,
1123
- "eval_samples_per_second": 579.815,
1124
- "eval_steps_per_second": 10.232,
1125
  "step": 3456
1126
  },
1127
  {
@@ -1150,9 +1150,9 @@
1150
  "eval_overall_f1": 0.9316455696202531,
1151
  "eval_overall_precision": 0.9246231155778895,
1152
  "eval_overall_recall": 0.9387755102040817,
1153
- "eval_runtime": 0.2946,
1154
- "eval_samples_per_second": 577.02,
1155
- "eval_steps_per_second": 10.183,
1156
  "step": 3552
1157
  },
1158
  {
@@ -1181,9 +1181,9 @@
1181
  "eval_overall_f1": 0.9341772151898734,
1182
  "eval_overall_precision": 0.9271356783919598,
1183
  "eval_overall_recall": 0.9413265306122449,
1184
- "eval_runtime": 0.2932,
1185
- "eval_samples_per_second": 579.802,
1186
- "eval_steps_per_second": 10.232,
1187
  "step": 3648
1188
  },
1189
  {
@@ -1212,9 +1212,9 @@
1212
  "eval_overall_f1": 0.9378960709759189,
1213
  "eval_overall_precision": 0.9319899244332494,
1214
  "eval_overall_recall": 0.9438775510204082,
1215
- "eval_runtime": 0.2963,
1216
- "eval_samples_per_second": 573.831,
1217
- "eval_steps_per_second": 10.126,
1218
  "step": 3744
1219
  },
1220
  {
@@ -1243,9 +1243,9 @@
1243
  "eval_overall_f1": 0.9428208386277002,
1244
  "eval_overall_precision": 0.9392405063291139,
1245
  "eval_overall_recall": 0.9464285714285714,
1246
- "eval_runtime": 0.2936,
1247
- "eval_samples_per_second": 579.089,
1248
- "eval_steps_per_second": 10.219,
1249
  "step": 3840
1250
  },
1251
  {
@@ -1274,9 +1274,9 @@
1274
  "eval_overall_f1": 0.9311224489795918,
1275
  "eval_overall_precision": 0.9311224489795918,
1276
  "eval_overall_recall": 0.9311224489795918,
1277
- "eval_runtime": 0.2935,
1278
- "eval_samples_per_second": 579.212,
1279
- "eval_steps_per_second": 10.221,
1280
  "step": 3936
1281
  },
1282
  {
@@ -1305,9 +1305,9 @@
1305
  "eval_overall_f1": 0.9318181818181819,
1306
  "eval_overall_precision": 0.9225,
1307
  "eval_overall_recall": 0.9413265306122449,
1308
- "eval_runtime": 0.2939,
1309
- "eval_samples_per_second": 578.465,
1310
- "eval_steps_per_second": 10.208,
1311
  "step": 4032
1312
  },
1313
  {
@@ -1336,9 +1336,9 @@
1336
  "eval_overall_f1": 0.935361216730038,
1337
  "eval_overall_precision": 0.929471032745592,
1338
  "eval_overall_recall": 0.9413265306122449,
1339
- "eval_runtime": 0.2928,
1340
- "eval_samples_per_second": 580.61,
1341
- "eval_steps_per_second": 10.246,
1342
  "step": 4128
1343
  },
1344
  {
@@ -1367,9 +1367,9 @@
1367
  "eval_overall_f1": 0.9365482233502538,
1368
  "eval_overall_precision": 0.9318181818181818,
1369
  "eval_overall_recall": 0.9413265306122449,
1370
- "eval_runtime": 0.2941,
1371
- "eval_samples_per_second": 578.039,
1372
- "eval_steps_per_second": 10.201,
1373
  "step": 4224
1374
  },
1375
  {
@@ -1398,9 +1398,9 @@
1398
  "eval_overall_f1": 0.9238578680203046,
1399
  "eval_overall_precision": 0.9191919191919192,
1400
  "eval_overall_recall": 0.9285714285714286,
1401
- "eval_runtime": 0.298,
1402
- "eval_samples_per_second": 570.445,
1403
- "eval_steps_per_second": 10.067,
1404
  "step": 4320
1405
  },
1406
  {
@@ -1429,9 +1429,9 @@
1429
  "eval_overall_f1": 0.9301143583227446,
1430
  "eval_overall_precision": 0.9265822784810127,
1431
  "eval_overall_recall": 0.9336734693877551,
1432
- "eval_runtime": 0.2941,
1433
- "eval_samples_per_second": 578.035,
1434
- "eval_steps_per_second": 10.201,
1435
  "step": 4416
1436
  },
1437
  {
@@ -1460,9 +1460,9 @@
1460
  "eval_overall_f1": 0.9367088607594937,
1461
  "eval_overall_precision": 0.9296482412060302,
1462
  "eval_overall_recall": 0.9438775510204082,
1463
- "eval_runtime": 0.293,
1464
- "eval_samples_per_second": 580.302,
1465
- "eval_steps_per_second": 10.241,
1466
  "step": 4512
1467
  },
1468
  {
@@ -1491,9 +1491,9 @@
1491
  "eval_overall_f1": 0.9316455696202531,
1492
  "eval_overall_precision": 0.9246231155778895,
1493
  "eval_overall_recall": 0.9387755102040817,
1494
- "eval_runtime": 0.2969,
1495
- "eval_samples_per_second": 572.674,
1496
- "eval_steps_per_second": 10.106,
1497
  "step": 4608
1498
  },
1499
  {
@@ -1522,9 +1522,9 @@
1522
  "eval_overall_f1": 0.937579617834395,
1523
  "eval_overall_precision": 0.9363867684478372,
1524
  "eval_overall_recall": 0.9387755102040817,
1525
- "eval_runtime": 0.2994,
1526
- "eval_samples_per_second": 567.886,
1527
- "eval_steps_per_second": 10.022,
1528
  "step": 4704
1529
  },
1530
  {
@@ -1553,9 +1553,9 @@
1553
  "eval_overall_f1": 0.9480354879594423,
1554
  "eval_overall_precision": 0.9420654911838791,
1555
  "eval_overall_recall": 0.9540816326530612,
1556
- "eval_runtime": 0.3003,
1557
- "eval_samples_per_second": 566.076,
1558
- "eval_steps_per_second": 9.99,
1559
  "step": 4800
1560
  },
1561
  {
@@ -1584,9 +1584,9 @@
1584
  "eval_overall_f1": 0.94147582697201,
1585
  "eval_overall_precision": 0.9390862944162437,
1586
  "eval_overall_recall": 0.9438775510204082,
1587
- "eval_runtime": 0.2998,
1588
- "eval_samples_per_second": 567.025,
1589
- "eval_steps_per_second": 10.006,
1590
  "step": 4896
1591
  },
1592
  {
@@ -1615,9 +1615,9 @@
1615
  "eval_overall_f1": 0.9402795425667091,
1616
  "eval_overall_precision": 0.9367088607594937,
1617
  "eval_overall_recall": 0.9438775510204082,
1618
- "eval_runtime": 0.2995,
1619
- "eval_samples_per_second": 567.671,
1620
- "eval_steps_per_second": 10.018,
1621
  "step": 4992
1622
  },
1623
  {
@@ -1646,9 +1646,9 @@
1646
  "eval_overall_f1": 0.940127388535032,
1647
  "eval_overall_precision": 0.9389312977099237,
1648
  "eval_overall_recall": 0.9413265306122449,
1649
- "eval_runtime": 0.2965,
1650
- "eval_samples_per_second": 573.314,
1651
- "eval_steps_per_second": 10.117,
1652
  "step": 5088
1653
  },
1654
  {
@@ -1677,9 +1677,9 @@
1677
  "eval_overall_f1": 0.9326556543837357,
1678
  "eval_overall_precision": 0.9291139240506329,
1679
  "eval_overall_recall": 0.9362244897959183,
1680
- "eval_runtime": 0.2958,
1681
- "eval_samples_per_second": 574.709,
1682
- "eval_steps_per_second": 10.142,
1683
  "step": 5184
1684
  },
1685
  {
@@ -1708,9 +1708,9 @@
1708
  "eval_overall_f1": 0.9411764705882353,
1709
  "eval_overall_precision": 0.9435897435897436,
1710
  "eval_overall_recall": 0.9387755102040817,
1711
- "eval_runtime": 0.2958,
1712
- "eval_samples_per_second": 574.619,
1713
- "eval_steps_per_second": 10.14,
1714
  "step": 5280
1715
  },
1716
  {
@@ -1739,9 +1739,9 @@
1739
  "eval_overall_f1": 0.9404309252217997,
1740
  "eval_overall_precision": 0.9345088161209067,
1741
  "eval_overall_recall": 0.9464285714285714,
1742
- "eval_runtime": 0.2983,
1743
- "eval_samples_per_second": 569.801,
1744
- "eval_steps_per_second": 10.055,
1745
  "step": 5376
1746
  },
1747
  {
@@ -1770,9 +1770,9 @@
1770
  "eval_overall_f1": 0.9417721518987343,
1771
  "eval_overall_precision": 0.9346733668341709,
1772
  "eval_overall_recall": 0.9489795918367347,
1773
- "eval_runtime": 0.2983,
1774
- "eval_samples_per_second": 569.857,
1775
- "eval_steps_per_second": 10.056,
1776
  "step": 5472
1777
  },
1778
  {
@@ -1801,9 +1801,9 @@
1801
  "eval_overall_f1": 0.9440203562340966,
1802
  "eval_overall_precision": 0.9416243654822335,
1803
  "eval_overall_recall": 0.9464285714285714,
1804
- "eval_runtime": 0.2981,
1805
- "eval_samples_per_second": 570.359,
1806
- "eval_steps_per_second": 10.065,
1807
  "step": 5568
1808
  },
1809
  {
@@ -1832,9 +1832,9 @@
1832
  "eval_overall_f1": 0.9402795425667091,
1833
  "eval_overall_precision": 0.9367088607594937,
1834
  "eval_overall_recall": 0.9438775510204082,
1835
- "eval_runtime": 0.2991,
1836
- "eval_samples_per_second": 568.459,
1837
- "eval_steps_per_second": 10.032,
1838
  "step": 5664
1839
  },
1840
  {
@@ -1863,9 +1863,9 @@
1863
  "eval_overall_f1": 0.9453621346886911,
1864
  "eval_overall_precision": 0.9417721518987342,
1865
  "eval_overall_recall": 0.9489795918367347,
1866
- "eval_runtime": 0.3001,
1867
- "eval_samples_per_second": 566.448,
1868
- "eval_steps_per_second": 9.996,
1869
  "step": 5760
1870
  },
1871
  {
@@ -1894,9 +1894,9 @@
1894
  "eval_overall_f1": 0.9426751592356687,
1895
  "eval_overall_precision": 0.9414758269720102,
1896
  "eval_overall_recall": 0.9438775510204082,
1897
- "eval_runtime": 0.3031,
1898
- "eval_samples_per_second": 560.797,
1899
- "eval_steps_per_second": 9.896,
1900
  "step": 5856
1901
  },
1902
  {
@@ -1925,9 +1925,9 @@
1925
  "eval_overall_f1": 0.9328263624841572,
1926
  "eval_overall_precision": 0.9269521410579346,
1927
  "eval_overall_recall": 0.9387755102040817,
1928
- "eval_runtime": 0.3001,
1929
- "eval_samples_per_second": 566.466,
1930
- "eval_steps_per_second": 9.996,
1931
  "step": 5952
1932
  },
1933
  {
@@ -1956,9 +1956,9 @@
1956
  "eval_overall_f1": 0.9377382465057178,
1957
  "eval_overall_precision": 0.9341772151898734,
1958
  "eval_overall_recall": 0.9413265306122449,
1959
- "eval_runtime": 0.3007,
1960
- "eval_samples_per_second": 565.435,
1961
- "eval_steps_per_second": 9.978,
1962
  "step": 6048
1963
  },
1964
  {
@@ -1987,9 +1987,9 @@
1987
  "eval_overall_f1": 0.9365482233502538,
1988
  "eval_overall_precision": 0.9318181818181818,
1989
  "eval_overall_recall": 0.9413265306122449,
1990
- "eval_runtime": 0.2971,
1991
- "eval_samples_per_second": 572.139,
1992
- "eval_steps_per_second": 10.097,
1993
  "step": 6144
1994
  },
1995
  {
@@ -2018,9 +2018,9 @@
2018
  "eval_overall_f1": 0.9405815423514539,
2019
  "eval_overall_precision": 0.9323308270676691,
2020
  "eval_overall_recall": 0.9489795918367347,
2021
- "eval_runtime": 0.2954,
2022
- "eval_samples_per_second": 575.53,
2023
- "eval_steps_per_second": 10.156,
2024
  "step": 6240
2025
  },
2026
  {
@@ -2049,9 +2049,9 @@
2049
  "eval_overall_f1": 0.94147582697201,
2050
  "eval_overall_precision": 0.9390862944162437,
2051
  "eval_overall_recall": 0.9438775510204082,
2052
- "eval_runtime": 0.2962,
2053
- "eval_samples_per_second": 573.923,
2054
- "eval_steps_per_second": 10.128,
2055
  "step": 6336
2056
  },
2057
  {
@@ -2080,9 +2080,9 @@
2080
  "eval_overall_f1": 0.9377382465057178,
2081
  "eval_overall_precision": 0.9341772151898734,
2082
  "eval_overall_recall": 0.9413265306122449,
2083
- "eval_runtime": 0.2956,
2084
- "eval_samples_per_second": 575.147,
2085
- "eval_steps_per_second": 10.15,
2086
  "step": 6432
2087
  },
2088
  {
@@ -2111,9 +2111,9 @@
2111
  "eval_overall_f1": 0.9377382465057178,
2112
  "eval_overall_precision": 0.9341772151898734,
2113
  "eval_overall_recall": 0.9413265306122449,
2114
- "eval_runtime": 0.3073,
2115
- "eval_samples_per_second": 553.256,
2116
- "eval_steps_per_second": 9.763,
2117
  "step": 6528
2118
  },
2119
  {
@@ -2142,9 +2142,9 @@
2142
  "eval_overall_f1": 0.9402795425667091,
2143
  "eval_overall_precision": 0.9367088607594937,
2144
  "eval_overall_recall": 0.9438775510204082,
2145
- "eval_runtime": 0.2944,
2146
- "eval_samples_per_second": 577.35,
2147
- "eval_steps_per_second": 10.189,
2148
  "step": 6624
2149
  },
2150
  {
@@ -2173,9 +2173,9 @@
2173
  "eval_overall_f1": 0.9377382465057178,
2174
  "eval_overall_precision": 0.9341772151898734,
2175
  "eval_overall_recall": 0.9413265306122449,
2176
- "eval_runtime": 0.2952,
2177
- "eval_samples_per_second": 575.901,
2178
- "eval_steps_per_second": 10.163,
2179
  "step": 6720
2180
  },
2181
  {
@@ -2204,9 +2204,9 @@
2204
  "eval_overall_f1": 0.9479034307496824,
2205
  "eval_overall_precision": 0.9443037974683545,
2206
  "eval_overall_recall": 0.951530612244898,
2207
- "eval_runtime": 0.2941,
2208
- "eval_samples_per_second": 578.109,
2209
- "eval_steps_per_second": 10.202,
2210
  "step": 6816
2211
  },
2212
  {
@@ -2235,9 +2235,9 @@
2235
  "eval_overall_f1": 0.9425287356321839,
2236
  "eval_overall_precision": 0.9437340153452686,
2237
  "eval_overall_recall": 0.9413265306122449,
2238
- "eval_runtime": 0.2935,
2239
- "eval_samples_per_second": 579.119,
2240
- "eval_steps_per_second": 10.22,
2241
  "step": 6912
2242
  },
2243
  {
@@ -2266,9 +2266,9 @@
2266
  "eval_overall_f1": 0.940127388535032,
2267
  "eval_overall_precision": 0.9389312977099237,
2268
  "eval_overall_recall": 0.9413265306122449,
2269
- "eval_runtime": 0.2989,
2270
- "eval_samples_per_second": 568.786,
2271
- "eval_steps_per_second": 10.037,
2272
  "step": 7008
2273
  },
2274
  {
@@ -2297,9 +2297,9 @@
2297
  "eval_overall_f1": 0.9413265306122449,
2298
  "eval_overall_precision": 0.9413265306122449,
2299
  "eval_overall_recall": 0.9413265306122449,
2300
- "eval_runtime": 0.3047,
2301
- "eval_samples_per_second": 558.015,
2302
- "eval_steps_per_second": 9.847,
2303
  "step": 7104
2304
  },
2305
  {
@@ -2328,9 +2328,9 @@
2328
  "eval_overall_f1": 0.9363867684478372,
2329
  "eval_overall_precision": 0.934010152284264,
2330
  "eval_overall_recall": 0.9387755102040817,
2331
- "eval_runtime": 0.2968,
2332
- "eval_samples_per_second": 572.822,
2333
- "eval_steps_per_second": 10.109,
2334
  "step": 7200
2335
  },
2336
  {
@@ -2359,9 +2359,9 @@
2359
  "eval_overall_f1": 0.9389312977099236,
2360
  "eval_overall_precision": 0.9365482233502538,
2361
  "eval_overall_recall": 0.9413265306122449,
2362
- "eval_runtime": 0.2937,
2363
- "eval_samples_per_second": 578.767,
2364
- "eval_steps_per_second": 10.214,
2365
  "step": 7296
2366
  },
2367
  {
@@ -2390,9 +2390,9 @@
2390
  "eval_overall_f1": 0.9387755102040817,
2391
  "eval_overall_precision": 0.9387755102040817,
2392
  "eval_overall_recall": 0.9387755102040817,
2393
- "eval_runtime": 0.2931,
2394
- "eval_samples_per_second": 579.925,
2395
- "eval_steps_per_second": 10.234,
2396
  "step": 7392
2397
  },
2398
  {
@@ -2421,9 +2421,9 @@
2421
  "eval_overall_f1": 0.9435897435897437,
2422
  "eval_overall_precision": 0.9484536082474226,
2423
  "eval_overall_recall": 0.9387755102040817,
2424
- "eval_runtime": 0.294,
2425
- "eval_samples_per_second": 578.15,
2426
- "eval_steps_per_second": 10.203,
2427
  "step": 7488
2428
  },
2429
  {
@@ -2452,9 +2452,9 @@
2452
  "eval_overall_f1": 0.9389312977099236,
2453
  "eval_overall_precision": 0.9365482233502538,
2454
  "eval_overall_recall": 0.9413265306122449,
2455
- "eval_runtime": 0.2953,
2456
- "eval_samples_per_second": 575.777,
2457
- "eval_steps_per_second": 10.161,
2458
  "step": 7584
2459
  },
2460
  {
@@ -2483,9 +2483,9 @@
2483
  "eval_overall_f1": 0.9453621346886911,
2484
  "eval_overall_precision": 0.9417721518987342,
2485
  "eval_overall_recall": 0.9489795918367347,
2486
- "eval_runtime": 0.2931,
2487
- "eval_samples_per_second": 579.941,
2488
- "eval_steps_per_second": 10.234,
2489
  "step": 7680
2490
  },
2491
  {
@@ -2514,9 +2514,9 @@
2514
  "eval_overall_f1": 0.937579617834395,
2515
  "eval_overall_precision": 0.9363867684478372,
2516
  "eval_overall_recall": 0.9387755102040817,
2517
- "eval_runtime": 0.2966,
2518
- "eval_samples_per_second": 573.117,
2519
- "eval_steps_per_second": 10.114,
2520
  "step": 7776
2521
  },
2522
  {
@@ -2545,9 +2545,9 @@
2545
  "eval_overall_f1": 0.9413265306122449,
2546
  "eval_overall_precision": 0.9413265306122449,
2547
  "eval_overall_recall": 0.9413265306122449,
2548
- "eval_runtime": 0.2951,
2549
- "eval_samples_per_second": 576.001,
2550
- "eval_steps_per_second": 10.165,
2551
  "step": 7872
2552
  },
2553
  {
@@ -2576,9 +2576,9 @@
2576
  "eval_overall_f1": 0.940127388535032,
2577
  "eval_overall_precision": 0.9389312977099237,
2578
  "eval_overall_recall": 0.9413265306122449,
2579
- "eval_runtime": 0.2928,
2580
- "eval_samples_per_second": 580.672,
2581
- "eval_steps_per_second": 10.247,
2582
  "step": 7968
2583
  },
2584
  {
@@ -2607,9 +2607,9 @@
2607
  "eval_overall_f1": 0.9390862944162437,
2608
  "eval_overall_precision": 0.9343434343434344,
2609
  "eval_overall_recall": 0.9438775510204082,
2610
- "eval_runtime": 0.2952,
2611
- "eval_samples_per_second": 575.828,
2612
- "eval_steps_per_second": 10.162,
2613
  "step": 8064
2614
  },
2615
  {
@@ -2638,9 +2638,9 @@
2638
  "eval_overall_f1": 0.9441624365482234,
2639
  "eval_overall_precision": 0.9393939393939394,
2640
  "eval_overall_recall": 0.9489795918367347,
2641
- "eval_runtime": 0.2934,
2642
- "eval_samples_per_second": 579.459,
2643
- "eval_steps_per_second": 10.226,
2644
  "step": 8160
2645
  },
2646
  {
@@ -2669,9 +2669,9 @@
2669
  "eval_overall_f1": 0.9402795425667091,
2670
  "eval_overall_precision": 0.9367088607594937,
2671
  "eval_overall_recall": 0.9438775510204082,
2672
- "eval_runtime": 0.2958,
2673
- "eval_samples_per_second": 574.664,
2674
- "eval_steps_per_second": 10.141,
2675
  "step": 8256
2676
  },
2677
  {
@@ -2700,9 +2700,9 @@
2700
  "eval_overall_f1": 0.9377382465057178,
2701
  "eval_overall_precision": 0.9341772151898734,
2702
  "eval_overall_recall": 0.9413265306122449,
2703
- "eval_runtime": 0.2941,
2704
- "eval_samples_per_second": 578.021,
2705
- "eval_steps_per_second": 10.2,
2706
  "step": 8352
2707
  },
2708
  {
@@ -2731,9 +2731,9 @@
2731
  "eval_overall_f1": 0.935361216730038,
2732
  "eval_overall_precision": 0.929471032745592,
2733
  "eval_overall_recall": 0.9413265306122449,
2734
- "eval_runtime": 0.2936,
2735
- "eval_samples_per_second": 579.002,
2736
- "eval_steps_per_second": 10.218,
2737
  "step": 8448
2738
  },
2739
  {
@@ -2762,9 +2762,9 @@
2762
  "eval_overall_f1": 0.9438775510204082,
2763
  "eval_overall_precision": 0.9438775510204082,
2764
  "eval_overall_recall": 0.9438775510204082,
2765
- "eval_runtime": 0.2938,
2766
- "eval_samples_per_second": 578.682,
2767
- "eval_steps_per_second": 10.212,
2768
  "step": 8544
2769
  },
2770
  {
@@ -2793,9 +2793,9 @@
2793
  "eval_overall_f1": 0.9314720812182741,
2794
  "eval_overall_precision": 0.9267676767676768,
2795
  "eval_overall_recall": 0.9362244897959183,
2796
- "eval_runtime": 0.2934,
2797
- "eval_samples_per_second": 579.347,
2798
- "eval_steps_per_second": 10.224,
2799
  "step": 8640
2800
  },
2801
  {
@@ -2824,9 +2824,9 @@
2824
  "eval_overall_f1": 0.9404309252217997,
2825
  "eval_overall_precision": 0.9345088161209067,
2826
  "eval_overall_recall": 0.9464285714285714,
2827
- "eval_runtime": 0.2921,
2828
- "eval_samples_per_second": 581.928,
2829
- "eval_steps_per_second": 10.269,
2830
  "step": 8736
2831
  },
2832
  {
@@ -2855,9 +2855,9 @@
2855
  "eval_overall_f1": 0.935361216730038,
2856
  "eval_overall_precision": 0.929471032745592,
2857
  "eval_overall_recall": 0.9413265306122449,
2858
- "eval_runtime": 0.294,
2859
- "eval_samples_per_second": 578.261,
2860
- "eval_steps_per_second": 10.205,
2861
  "step": 8832
2862
  },
2863
  {
@@ -2886,9 +2886,9 @@
2886
  "eval_overall_f1": 0.935361216730038,
2887
  "eval_overall_precision": 0.929471032745592,
2888
  "eval_overall_recall": 0.9413265306122449,
2889
- "eval_runtime": 0.295,
2890
- "eval_samples_per_second": 576.282,
2891
- "eval_steps_per_second": 10.17,
2892
  "step": 8928
2893
  },
2894
  {
@@ -2917,9 +2917,9 @@
2917
  "eval_overall_f1": 0.9416243654822335,
2918
  "eval_overall_precision": 0.9368686868686869,
2919
  "eval_overall_recall": 0.9464285714285714,
2920
- "eval_runtime": 0.2942,
2921
- "eval_samples_per_second": 577.791,
2922
- "eval_steps_per_second": 10.196,
2923
  "step": 9024
2924
  },
2925
  {
@@ -2948,9 +2948,9 @@
2948
  "eval_overall_f1": 0.9378960709759189,
2949
  "eval_overall_precision": 0.9319899244332494,
2950
  "eval_overall_recall": 0.9438775510204082,
2951
- "eval_runtime": 0.2953,
2952
- "eval_samples_per_second": 575.721,
2953
- "eval_steps_per_second": 10.16,
2954
  "step": 9120
2955
  },
2956
  {
@@ -2979,9 +2979,9 @@
2979
  "eval_overall_f1": 0.9341772151898734,
2980
  "eval_overall_precision": 0.9271356783919598,
2981
  "eval_overall_recall": 0.9413265306122449,
2982
- "eval_runtime": 0.2935,
2983
- "eval_samples_per_second": 579.192,
2984
- "eval_steps_per_second": 10.221,
2985
  "step": 9216
2986
  },
2987
  {
@@ -3010,9 +3010,9 @@
3010
  "eval_overall_f1": 0.9416243654822335,
3011
  "eval_overall_precision": 0.9368686868686869,
3012
  "eval_overall_recall": 0.9464285714285714,
3013
- "eval_runtime": 0.294,
3014
- "eval_samples_per_second": 578.159,
3015
- "eval_steps_per_second": 10.203,
3016
  "step": 9312
3017
  },
3018
  {
@@ -3041,9 +3041,9 @@
3041
  "eval_overall_f1": 0.9390862944162437,
3042
  "eval_overall_precision": 0.9343434343434344,
3043
  "eval_overall_recall": 0.9438775510204082,
3044
- "eval_runtime": 0.294,
3045
- "eval_samples_per_second": 578.285,
3046
- "eval_steps_per_second": 10.205,
3047
  "step": 9408
3048
  },
3049
  {
@@ -3072,9 +3072,9 @@
3072
  "eval_overall_f1": 0.9390862944162437,
3073
  "eval_overall_precision": 0.9343434343434344,
3074
  "eval_overall_recall": 0.9438775510204082,
3075
- "eval_runtime": 0.2926,
3076
- "eval_samples_per_second": 581.079,
3077
- "eval_steps_per_second": 10.254,
3078
  "step": 9504
3079
  },
3080
  {
@@ -3103,9 +3103,9 @@
3103
  "eval_overall_f1": 0.9390862944162437,
3104
  "eval_overall_precision": 0.9343434343434344,
3105
  "eval_overall_recall": 0.9438775510204082,
3106
- "eval_runtime": 0.2941,
3107
- "eval_samples_per_second": 578.016,
3108
- "eval_steps_per_second": 10.2,
3109
  "step": 9600
3110
  },
3111
  {
@@ -3113,9 +3113,9 @@
3113
  "step": 9600,
3114
  "total_flos": 4315798421360676.0,
3115
  "train_loss": 0.03753832100580136,
3116
- "train_runtime": 561.7168,
3117
- "train_samples_per_second": 272.557,
3118
- "train_steps_per_second": 17.09
3119
  }
3120
  ],
3121
  "logging_steps": 500,
 
34
  "eval_overall_f1": 0.23699421965317918,
35
  "eval_overall_precision": 0.2733333333333333,
36
  "eval_overall_recall": 0.20918367346938777,
37
+ "eval_runtime": 0.2945,
38
+ "eval_samples_per_second": 577.299,
39
+ "eval_steps_per_second": 10.188,
40
  "step": 96
41
  },
42
  {
 
65
  "eval_overall_f1": 0.585305105853051,
66
  "eval_overall_precision": 0.5717761557177615,
67
  "eval_overall_recall": 0.5994897959183674,
68
+ "eval_runtime": 0.3009,
69
+ "eval_samples_per_second": 564.972,
70
+ "eval_steps_per_second": 9.97,
71
  "step": 192
72
  },
73
  {
 
96
  "eval_overall_f1": 0.8114143920595533,
97
  "eval_overall_precision": 0.7898550724637681,
98
  "eval_overall_recall": 0.8341836734693877,
99
+ "eval_runtime": 0.2989,
100
+ "eval_samples_per_second": 568.667,
101
+ "eval_steps_per_second": 10.035,
102
  "step": 288
103
  },
104
  {
 
127
  "eval_overall_f1": 0.8335388409371147,
128
  "eval_overall_precision": 0.8066825775656324,
129
  "eval_overall_recall": 0.8622448979591837,
130
+ "eval_runtime": 0.3022,
131
+ "eval_samples_per_second": 562.561,
132
+ "eval_steps_per_second": 9.928,
133
  "step": 384
134
  },
135
  {
 
158
  "eval_overall_f1": 0.8955223880597014,
159
  "eval_overall_precision": 0.8737864077669902,
160
  "eval_overall_recall": 0.9183673469387755,
161
+ "eval_runtime": 0.3033,
162
+ "eval_samples_per_second": 560.446,
163
+ "eval_steps_per_second": 9.89,
164
  "step": 480
165
  },
166
  {
 
189
  "eval_overall_f1": 0.8664987405541562,
190
  "eval_overall_precision": 0.8557213930348259,
191
  "eval_overall_recall": 0.8775510204081632,
192
+ "eval_runtime": 0.2981,
193
+ "eval_samples_per_second": 570.228,
194
+ "eval_steps_per_second": 10.063,
195
  "step": 576
196
  },
197
  {
 
220
  "eval_overall_f1": 0.9052369077306733,
221
  "eval_overall_precision": 0.8853658536585366,
222
  "eval_overall_recall": 0.9260204081632653,
223
+ "eval_runtime": 0.2962,
224
+ "eval_samples_per_second": 573.963,
225
+ "eval_steps_per_second": 10.129,
226
  "step": 672
227
  },
228
  {
 
251
  "eval_overall_f1": 0.9081761006289308,
252
  "eval_overall_precision": 0.8957816377171216,
253
  "eval_overall_recall": 0.9209183673469388,
254
+ "eval_runtime": 0.2976,
255
+ "eval_samples_per_second": 571.277,
256
+ "eval_steps_per_second": 10.081,
257
  "step": 768
258
  },
259
  {
 
282
  "eval_overall_f1": 0.9051833122629582,
283
  "eval_overall_precision": 0.8972431077694235,
284
  "eval_overall_recall": 0.9132653061224489,
285
+ "eval_runtime": 0.3,
286
+ "eval_samples_per_second": 566.755,
287
+ "eval_steps_per_second": 10.002,
288
  "step": 864
289
  },
290
  {
 
313
  "eval_overall_f1": 0.9120603015075376,
314
  "eval_overall_precision": 0.8985148514851485,
315
  "eval_overall_recall": 0.9260204081632653,
316
+ "eval_runtime": 0.3007,
317
+ "eval_samples_per_second": 565.294,
318
+ "eval_steps_per_second": 9.976,
319
  "step": 960
320
  },
321
  {
 
344
  "eval_overall_f1": 0.9345088161209069,
345
  "eval_overall_precision": 0.9228855721393034,
346
  "eval_overall_recall": 0.9464285714285714,
347
+ "eval_runtime": 0.2996,
348
+ "eval_samples_per_second": 567.344,
349
+ "eval_steps_per_second": 10.012,
350
  "step": 1056
351
  },
352
  {
 
375
  "eval_overall_f1": 0.9371859296482412,
376
  "eval_overall_precision": 0.9232673267326733,
377
  "eval_overall_recall": 0.951530612244898,
378
+ "eval_runtime": 0.2993,
379
+ "eval_samples_per_second": 567.954,
380
+ "eval_steps_per_second": 10.023,
381
  "step": 1152
382
  },
383
  {
 
406
  "eval_overall_f1": 0.9319899244332494,
407
  "eval_overall_precision": 0.9203980099502488,
408
  "eval_overall_recall": 0.9438775510204082,
409
+ "eval_runtime": 0.2985,
410
+ "eval_samples_per_second": 569.598,
411
+ "eval_steps_per_second": 10.052,
412
  "step": 1248
413
  },
414
  {
 
437
  "eval_overall_f1": 0.9362244897959183,
438
  "eval_overall_precision": 0.9362244897959183,
439
  "eval_overall_recall": 0.9362244897959183,
440
+ "eval_runtime": 0.2993,
441
+ "eval_samples_per_second": 567.929,
442
+ "eval_steps_per_second": 10.022,
443
  "step": 1344
444
  },
445
  {
 
468
  "eval_overall_f1": 0.9287531806615775,
469
  "eval_overall_precision": 0.9263959390862944,
470
  "eval_overall_recall": 0.9311224489795918,
471
+ "eval_runtime": 0.2999,
472
+ "eval_samples_per_second": 566.824,
473
+ "eval_steps_per_second": 10.003,
474
  "step": 1440
475
  },
476
  {
 
499
  "eval_overall_f1": 0.9457755359394704,
500
  "eval_overall_precision": 0.9351620947630923,
501
  "eval_overall_recall": 0.9566326530612245,
502
+ "eval_runtime": 0.2981,
503
+ "eval_samples_per_second": 570.24,
504
+ "eval_steps_per_second": 10.063,
505
  "step": 1536
506
  },
507
  {
 
530
  "eval_overall_f1": 0.9445843828715365,
531
  "eval_overall_precision": 0.9328358208955224,
532
  "eval_overall_recall": 0.9566326530612245,
533
+ "eval_runtime": 0.3012,
534
+ "eval_samples_per_second": 564.398,
535
+ "eval_steps_per_second": 9.96,
536
  "step": 1632
537
  },
538
  {
 
561
  "eval_overall_f1": 0.9489795918367347,
562
  "eval_overall_precision": 0.9489795918367347,
563
  "eval_overall_recall": 0.9489795918367347,
564
+ "eval_runtime": 0.2986,
565
+ "eval_samples_per_second": 569.248,
566
+ "eval_steps_per_second": 10.046,
567
  "step": 1728
568
  },
569
  {
 
592
  "eval_overall_f1": 0.9516539440203563,
593
  "eval_overall_precision": 0.949238578680203,
594
  "eval_overall_recall": 0.9540816326530612,
595
+ "eval_runtime": 0.2977,
596
+ "eval_samples_per_second": 571.08,
597
+ "eval_steps_per_second": 10.078,
598
  "step": 1824
599
  },
600
  {
 
623
  "eval_overall_f1": 0.9333333333333335,
624
  "eval_overall_precision": 0.9205955334987593,
625
  "eval_overall_recall": 0.9464285714285714,
626
+ "eval_runtime": 0.2974,
627
+ "eval_samples_per_second": 571.576,
628
+ "eval_steps_per_second": 10.087,
629
  "step": 1920
630
  },
631
  {
 
654
  "eval_overall_f1": 0.9402795425667091,
655
  "eval_overall_precision": 0.9367088607594937,
656
  "eval_overall_recall": 0.9438775510204082,
657
+ "eval_runtime": 0.2981,
658
+ "eval_samples_per_second": 570.287,
659
+ "eval_steps_per_second": 10.064,
660
  "step": 2016
661
  },
662
  {
 
685
  "eval_overall_f1": 0.9213483146067415,
686
  "eval_overall_precision": 0.902200488997555,
687
  "eval_overall_recall": 0.9413265306122449,
688
+ "eval_runtime": 0.3002,
689
+ "eval_samples_per_second": 566.326,
690
+ "eval_steps_per_second": 9.994,
691
  "step": 2112
692
  },
693
  {
 
716
  "eval_overall_f1": 0.9444444444444445,
717
  "eval_overall_precision": 0.935,
718
  "eval_overall_recall": 0.9540816326530612,
719
+ "eval_runtime": 0.3015,
720
+ "eval_samples_per_second": 563.806,
721
+ "eval_steps_per_second": 9.95,
722
  "step": 2208
723
  },
724
  {
 
747
  "eval_overall_f1": 0.929113924050633,
748
  "eval_overall_precision": 0.9221105527638191,
749
  "eval_overall_recall": 0.9362244897959183,
750
+ "eval_runtime": 0.3025,
751
+ "eval_samples_per_second": 561.908,
752
+ "eval_steps_per_second": 9.916,
753
  "step": 2304
754
  },
755
  {
 
778
  "eval_overall_f1": 0.9438775510204082,
779
  "eval_overall_precision": 0.9438775510204082,
780
  "eval_overall_recall": 0.9438775510204082,
781
+ "eval_runtime": 0.3,
782
+ "eval_samples_per_second": 566.687,
783
+ "eval_steps_per_second": 10.0,
784
  "step": 2400
785
  },
786
  {
 
809
  "eval_overall_f1": 0.9390862944162437,
810
  "eval_overall_precision": 0.9343434343434344,
811
  "eval_overall_recall": 0.9438775510204082,
812
+ "eval_runtime": 0.2974,
813
+ "eval_samples_per_second": 571.593,
814
+ "eval_steps_per_second": 10.087,
815
  "step": 2496
816
  },
817
  {
 
840
  "eval_overall_f1": 0.9385194479297364,
841
  "eval_overall_precision": 0.9234567901234568,
842
  "eval_overall_recall": 0.9540816326530612,
843
+ "eval_runtime": 0.297,
844
+ "eval_samples_per_second": 572.323,
845
+ "eval_steps_per_second": 10.1,
846
  "step": 2592
847
  },
848
  {
 
871
  "eval_overall_f1": 0.9368686868686869,
872
  "eval_overall_precision": 0.9275,
873
  "eval_overall_recall": 0.9464285714285714,
874
+ "eval_runtime": 0.2987,
875
+ "eval_samples_per_second": 569.044,
876
+ "eval_steps_per_second": 10.042,
877
  "step": 2688
878
  },
879
  {
 
902
  "eval_overall_f1": 0.9343434343434343,
903
  "eval_overall_precision": 0.925,
904
  "eval_overall_recall": 0.9438775510204082,
905
+ "eval_runtime": 0.3063,
906
+ "eval_samples_per_second": 555.038,
907
+ "eval_steps_per_second": 9.795,
908
  "step": 2784
909
  },
910
  {
 
933
  "eval_overall_f1": 0.9433962264150944,
934
  "eval_overall_precision": 0.9305210918114144,
935
  "eval_overall_recall": 0.9566326530612245,
936
+ "eval_runtime": 0.3006,
937
+ "eval_samples_per_second": 565.457,
938
+ "eval_steps_per_second": 9.979,
939
  "step": 2880
940
  },
941
  {
 
964
  "eval_overall_f1": 0.9389312977099236,
965
  "eval_overall_precision": 0.9365482233502538,
966
  "eval_overall_recall": 0.9413265306122449,
967
+ "eval_runtime": 0.2984,
968
+ "eval_samples_per_second": 569.615,
969
+ "eval_steps_per_second": 10.052,
970
  "step": 2976
971
  },
972
  {
 
995
  "eval_overall_f1": 0.9328263624841572,
996
  "eval_overall_precision": 0.9269521410579346,
997
  "eval_overall_recall": 0.9387755102040817,
998
+ "eval_runtime": 0.3007,
999
+ "eval_samples_per_second": 565.425,
1000
+ "eval_steps_per_second": 9.978,
1001
  "step": 3072
1002
  },
1003
  {
 
1026
  "eval_overall_f1": 0.9491094147582698,
1027
  "eval_overall_precision": 0.9467005076142132,
1028
  "eval_overall_recall": 0.951530612244898,
1029
+ "eval_runtime": 0.3013,
1030
+ "eval_samples_per_second": 564.287,
1031
+ "eval_steps_per_second": 9.958,
1032
  "step": 3168
1033
  },
1034
  {
 
1057
  "eval_overall_f1": 0.9444444444444445,
1058
  "eval_overall_precision": 0.935,
1059
  "eval_overall_recall": 0.9540816326530612,
1060
+ "eval_runtime": 0.3009,
1061
+ "eval_samples_per_second": 564.913,
1062
+ "eval_steps_per_second": 9.969,
1063
  "step": 3264
1064
  },
1065
  {
 
1088
  "eval_overall_f1": 0.9429657794676806,
1089
  "eval_overall_precision": 0.9370277078085643,
1090
  "eval_overall_recall": 0.9489795918367347,
1091
+ "eval_runtime": 0.3012,
1092
+ "eval_samples_per_second": 564.434,
1093
+ "eval_steps_per_second": 9.961,
1094
  "step": 3360
1095
  },
1096
  {
 
1119
  "eval_overall_f1": 0.9211195928753181,
1120
  "eval_overall_precision": 0.9187817258883249,
1121
  "eval_overall_recall": 0.923469387755102,
1122
+ "eval_runtime": 0.2981,
1123
+ "eval_samples_per_second": 570.214,
1124
+ "eval_steps_per_second": 10.063,
1125
  "step": 3456
1126
  },
1127
  {
 
1150
  "eval_overall_f1": 0.9316455696202531,
1151
  "eval_overall_precision": 0.9246231155778895,
1152
  "eval_overall_recall": 0.9387755102040817,
1153
+ "eval_runtime": 0.3015,
1154
+ "eval_samples_per_second": 563.791,
1155
+ "eval_steps_per_second": 9.949,
1156
  "step": 3552
1157
  },
1158
  {
 
1181
  "eval_overall_f1": 0.9341772151898734,
1182
  "eval_overall_precision": 0.9271356783919598,
1183
  "eval_overall_recall": 0.9413265306122449,
1184
+ "eval_runtime": 0.2984,
1185
+ "eval_samples_per_second": 569.644,
1186
+ "eval_steps_per_second": 10.053,
1187
  "step": 3648
1188
  },
1189
  {
 
1212
  "eval_overall_f1": 0.9378960709759189,
1213
  "eval_overall_precision": 0.9319899244332494,
1214
  "eval_overall_recall": 0.9438775510204082,
1215
+ "eval_runtime": 0.2971,
1216
+ "eval_samples_per_second": 572.14,
1217
+ "eval_steps_per_second": 10.097,
1218
  "step": 3744
1219
  },
1220
  {
 
1243
  "eval_overall_f1": 0.9428208386277002,
1244
  "eval_overall_precision": 0.9392405063291139,
1245
  "eval_overall_recall": 0.9464285714285714,
1246
+ "eval_runtime": 0.3002,
1247
+ "eval_samples_per_second": 566.35,
1248
+ "eval_steps_per_second": 9.994,
1249
  "step": 3840
1250
  },
1251
  {
 
1274
  "eval_overall_f1": 0.9311224489795918,
1275
  "eval_overall_precision": 0.9311224489795918,
1276
  "eval_overall_recall": 0.9311224489795918,
1277
+ "eval_runtime": 0.2988,
1278
+ "eval_samples_per_second": 568.92,
1279
+ "eval_steps_per_second": 10.04,
1280
  "step": 3936
1281
  },
1282
  {
 
1305
  "eval_overall_f1": 0.9318181818181819,
1306
  "eval_overall_precision": 0.9225,
1307
  "eval_overall_recall": 0.9413265306122449,
1308
+ "eval_runtime": 0.2996,
1309
+ "eval_samples_per_second": 567.494,
1310
+ "eval_steps_per_second": 10.015,
1311
  "step": 4032
1312
  },
1313
  {
 
1336
  "eval_overall_f1": 0.935361216730038,
1337
  "eval_overall_precision": 0.929471032745592,
1338
  "eval_overall_recall": 0.9413265306122449,
1339
+ "eval_runtime": 0.2982,
1340
+ "eval_samples_per_second": 570.136,
1341
+ "eval_steps_per_second": 10.061,
1342
  "step": 4128
1343
  },
1344
  {
 
1367
  "eval_overall_f1": 0.9365482233502538,
1368
  "eval_overall_precision": 0.9318181818181818,
1369
  "eval_overall_recall": 0.9413265306122449,
1370
+ "eval_runtime": 0.2982,
1371
+ "eval_samples_per_second": 570.013,
1372
+ "eval_steps_per_second": 10.059,
1373
  "step": 4224
1374
  },
1375
  {
 
1398
  "eval_overall_f1": 0.9238578680203046,
1399
  "eval_overall_precision": 0.9191919191919192,
1400
  "eval_overall_recall": 0.9285714285714286,
1401
+ "eval_runtime": 0.2963,
1402
+ "eval_samples_per_second": 573.707,
1403
+ "eval_steps_per_second": 10.124,
1404
  "step": 4320
1405
  },
1406
  {
 
1429
  "eval_overall_f1": 0.9301143583227446,
1430
  "eval_overall_precision": 0.9265822784810127,
1431
  "eval_overall_recall": 0.9336734693877551,
1432
+ "eval_runtime": 0.2962,
1433
+ "eval_samples_per_second": 573.994,
1434
+ "eval_steps_per_second": 10.129,
1435
  "step": 4416
1436
  },
1437
  {
 
1460
  "eval_overall_f1": 0.9367088607594937,
1461
  "eval_overall_precision": 0.9296482412060302,
1462
  "eval_overall_recall": 0.9438775510204082,
1463
+ "eval_runtime": 0.2997,
1464
+ "eval_samples_per_second": 567.195,
1465
+ "eval_steps_per_second": 10.009,
1466
  "step": 4512
1467
  },
1468
  {
 
1491
  "eval_overall_f1": 0.9316455696202531,
1492
  "eval_overall_precision": 0.9246231155778895,
1493
  "eval_overall_recall": 0.9387755102040817,
1494
+ "eval_runtime": 0.2958,
1495
+ "eval_samples_per_second": 574.809,
1496
+ "eval_steps_per_second": 10.144,
1497
  "step": 4608
1498
  },
1499
  {
 
1522
  "eval_overall_f1": 0.937579617834395,
1523
  "eval_overall_precision": 0.9363867684478372,
1524
  "eval_overall_recall": 0.9387755102040817,
1525
+ "eval_runtime": 0.2963,
1526
+ "eval_samples_per_second": 573.646,
1527
+ "eval_steps_per_second": 10.123,
1528
  "step": 4704
1529
  },
1530
  {
 
1553
  "eval_overall_f1": 0.9480354879594423,
1554
  "eval_overall_precision": 0.9420654911838791,
1555
  "eval_overall_recall": 0.9540816326530612,
1556
+ "eval_runtime": 0.2947,
1557
+ "eval_samples_per_second": 576.898,
1558
+ "eval_steps_per_second": 10.181,
1559
  "step": 4800
1560
  },
1561
  {
 
1584
  "eval_overall_f1": 0.94147582697201,
1585
  "eval_overall_precision": 0.9390862944162437,
1586
  "eval_overall_recall": 0.9438775510204082,
1587
+ "eval_runtime": 0.294,
1588
+ "eval_samples_per_second": 578.322,
1589
+ "eval_steps_per_second": 10.206,
1590
  "step": 4896
1591
  },
1592
  {
 
1615
  "eval_overall_f1": 0.9402795425667091,
1616
  "eval_overall_precision": 0.9367088607594937,
1617
  "eval_overall_recall": 0.9438775510204082,
1618
+ "eval_runtime": 0.2951,
1619
+ "eval_samples_per_second": 576.102,
1620
+ "eval_steps_per_second": 10.167,
1621
  "step": 4992
1622
  },
1623
  {
 
1646
  "eval_overall_f1": 0.940127388535032,
1647
  "eval_overall_precision": 0.9389312977099237,
1648
  "eval_overall_recall": 0.9413265306122449,
1649
+ "eval_runtime": 0.2937,
1650
+ "eval_samples_per_second": 578.915,
1651
+ "eval_steps_per_second": 10.216,
1652
  "step": 5088
1653
  },
1654
  {
 
1677
  "eval_overall_f1": 0.9326556543837357,
1678
  "eval_overall_precision": 0.9291139240506329,
1679
  "eval_overall_recall": 0.9362244897959183,
1680
+ "eval_runtime": 0.2955,
1681
+ "eval_samples_per_second": 575.218,
1682
+ "eval_steps_per_second": 10.151,
1683
  "step": 5184
1684
  },
1685
  {
 
1708
  "eval_overall_f1": 0.9411764705882353,
1709
  "eval_overall_precision": 0.9435897435897436,
1710
  "eval_overall_recall": 0.9387755102040817,
1711
+ "eval_runtime": 0.2961,
1712
+ "eval_samples_per_second": 574.105,
1713
+ "eval_steps_per_second": 10.131,
1714
  "step": 5280
1715
  },
1716
  {
 
1739
  "eval_overall_f1": 0.9404309252217997,
1740
  "eval_overall_precision": 0.9345088161209067,
1741
  "eval_overall_recall": 0.9464285714285714,
1742
+ "eval_runtime": 0.2987,
1743
+ "eval_samples_per_second": 569.153,
1744
+ "eval_steps_per_second": 10.044,
1745
  "step": 5376
1746
  },
1747
  {
 
1770
  "eval_overall_f1": 0.9417721518987343,
1771
  "eval_overall_precision": 0.9346733668341709,
1772
  "eval_overall_recall": 0.9489795918367347,
1773
+ "eval_runtime": 0.2951,
1774
+ "eval_samples_per_second": 576.154,
1775
+ "eval_steps_per_second": 10.167,
1776
  "step": 5472
1777
  },
1778
  {
 
1801
  "eval_overall_f1": 0.9440203562340966,
1802
  "eval_overall_precision": 0.9416243654822335,
1803
  "eval_overall_recall": 0.9464285714285714,
1804
+ "eval_runtime": 0.2948,
1805
+ "eval_samples_per_second": 576.685,
1806
+ "eval_steps_per_second": 10.177,
1807
  "step": 5568
1808
  },
1809
  {
 
1832
  "eval_overall_f1": 0.9402795425667091,
1833
  "eval_overall_precision": 0.9367088607594937,
1834
  "eval_overall_recall": 0.9438775510204082,
1835
+ "eval_runtime": 0.2959,
1836
+ "eval_samples_per_second": 574.543,
1837
+ "eval_steps_per_second": 10.139,
1838
  "step": 5664
1839
  },
1840
  {
 
1863
  "eval_overall_f1": 0.9453621346886911,
1864
  "eval_overall_precision": 0.9417721518987342,
1865
  "eval_overall_recall": 0.9489795918367347,
1866
+ "eval_runtime": 0.2966,
1867
+ "eval_samples_per_second": 573.16,
1868
+ "eval_steps_per_second": 10.115,
1869
  "step": 5760
1870
  },
1871
  {
 
1894
  "eval_overall_f1": 0.9426751592356687,
1895
  "eval_overall_precision": 0.9414758269720102,
1896
  "eval_overall_recall": 0.9438775510204082,
1897
+ "eval_runtime": 0.2935,
1898
+ "eval_samples_per_second": 579.297,
1899
+ "eval_steps_per_second": 10.223,
1900
  "step": 5856
1901
  },
1902
  {
 
1925
  "eval_overall_f1": 0.9328263624841572,
1926
  "eval_overall_precision": 0.9269521410579346,
1927
  "eval_overall_recall": 0.9387755102040817,
1928
+ "eval_runtime": 0.2962,
1929
+ "eval_samples_per_second": 573.927,
1930
+ "eval_steps_per_second": 10.128,
1931
  "step": 5952
1932
  },
1933
  {
 
1956
  "eval_overall_f1": 0.9377382465057178,
1957
  "eval_overall_precision": 0.9341772151898734,
1958
  "eval_overall_recall": 0.9413265306122449,
1959
+ "eval_runtime": 0.2978,
1960
+ "eval_samples_per_second": 570.869,
1961
+ "eval_steps_per_second": 10.074,
1962
  "step": 6048
1963
  },
1964
  {
 
1987
  "eval_overall_f1": 0.9365482233502538,
1988
  "eval_overall_precision": 0.9318181818181818,
1989
  "eval_overall_recall": 0.9413265306122449,
1990
+ "eval_runtime": 0.2951,
1991
+ "eval_samples_per_second": 576.125,
1992
+ "eval_steps_per_second": 10.167,
1993
  "step": 6144
1994
  },
1995
  {
 
2018
  "eval_overall_f1": 0.9405815423514539,
2019
  "eval_overall_precision": 0.9323308270676691,
2020
  "eval_overall_recall": 0.9489795918367347,
2021
+ "eval_runtime": 0.296,
2022
+ "eval_samples_per_second": 574.32,
2023
+ "eval_steps_per_second": 10.135,
2024
  "step": 6240
2025
  },
2026
  {
 
2049
  "eval_overall_f1": 0.94147582697201,
2050
  "eval_overall_precision": 0.9390862944162437,
2051
  "eval_overall_recall": 0.9438775510204082,
2052
+ "eval_runtime": 0.295,
2053
+ "eval_samples_per_second": 576.188,
2054
+ "eval_steps_per_second": 10.168,
2055
  "step": 6336
2056
  },
2057
  {
 
2080
  "eval_overall_f1": 0.9377382465057178,
2081
  "eval_overall_precision": 0.9341772151898734,
2082
  "eval_overall_recall": 0.9413265306122449,
2083
+ "eval_runtime": 0.2948,
2084
+ "eval_samples_per_second": 576.601,
2085
+ "eval_steps_per_second": 10.175,
2086
  "step": 6432
2087
  },
2088
  {
 
2111
  "eval_overall_f1": 0.9377382465057178,
2112
  "eval_overall_precision": 0.9341772151898734,
2113
  "eval_overall_recall": 0.9413265306122449,
2114
+ "eval_runtime": 0.2971,
2115
+ "eval_samples_per_second": 572.269,
2116
+ "eval_steps_per_second": 10.099,
2117
  "step": 6528
2118
  },
2119
  {
 
2142
  "eval_overall_f1": 0.9402795425667091,
2143
  "eval_overall_precision": 0.9367088607594937,
2144
  "eval_overall_recall": 0.9438775510204082,
2145
+ "eval_runtime": 0.2951,
2146
+ "eval_samples_per_second": 576.066,
2147
+ "eval_steps_per_second": 10.166,
2148
  "step": 6624
2149
  },
2150
  {
 
2173
  "eval_overall_f1": 0.9377382465057178,
2174
  "eval_overall_precision": 0.9341772151898734,
2175
  "eval_overall_recall": 0.9413265306122449,
2176
+ "eval_runtime": 0.2955,
2177
+ "eval_samples_per_second": 575.219,
2178
+ "eval_steps_per_second": 10.151,
2179
  "step": 6720
2180
  },
2181
  {
 
2204
  "eval_overall_f1": 0.9479034307496824,
2205
  "eval_overall_precision": 0.9443037974683545,
2206
  "eval_overall_recall": 0.951530612244898,
2207
+ "eval_runtime": 0.296,
2208
+ "eval_samples_per_second": 574.269,
2209
+ "eval_steps_per_second": 10.134,
2210
  "step": 6816
2211
  },
2212
  {
 
2235
  "eval_overall_f1": 0.9425287356321839,
2236
  "eval_overall_precision": 0.9437340153452686,
2237
  "eval_overall_recall": 0.9413265306122449,
2238
+ "eval_runtime": 0.2954,
2239
+ "eval_samples_per_second": 575.575,
2240
+ "eval_steps_per_second": 10.157,
2241
  "step": 6912
2242
  },
2243
  {
 
2266
  "eval_overall_f1": 0.940127388535032,
2267
  "eval_overall_precision": 0.9389312977099237,
2268
  "eval_overall_recall": 0.9413265306122449,
2269
+ "eval_runtime": 0.2935,
2270
+ "eval_samples_per_second": 579.206,
2271
+ "eval_steps_per_second": 10.221,
2272
  "step": 7008
2273
  },
2274
  {
 
2297
  "eval_overall_f1": 0.9413265306122449,
2298
  "eval_overall_precision": 0.9413265306122449,
2299
  "eval_overall_recall": 0.9413265306122449,
2300
+ "eval_runtime": 0.295,
2301
+ "eval_samples_per_second": 576.232,
2302
+ "eval_steps_per_second": 10.169,
2303
  "step": 7104
2304
  },
2305
  {
 
2328
  "eval_overall_f1": 0.9363867684478372,
2329
  "eval_overall_precision": 0.934010152284264,
2330
  "eval_overall_recall": 0.9387755102040817,
2331
+ "eval_runtime": 0.2925,
2332
+ "eval_samples_per_second": 581.225,
2333
+ "eval_steps_per_second": 10.257,
2334
  "step": 7200
2335
  },
2336
  {
 
2359
  "eval_overall_f1": 0.9389312977099236,
2360
  "eval_overall_precision": 0.9365482233502538,
2361
  "eval_overall_recall": 0.9413265306122449,
2362
+ "eval_runtime": 0.2954,
2363
+ "eval_samples_per_second": 575.462,
2364
+ "eval_steps_per_second": 10.155,
2365
  "step": 7296
2366
  },
2367
  {
 
2390
  "eval_overall_f1": 0.9387755102040817,
2391
  "eval_overall_precision": 0.9387755102040817,
2392
  "eval_overall_recall": 0.9387755102040817,
2393
+ "eval_runtime": 0.2953,
2394
+ "eval_samples_per_second": 575.722,
2395
+ "eval_steps_per_second": 10.16,
2396
  "step": 7392
2397
  },
2398
  {
 
2421
  "eval_overall_f1": 0.9435897435897437,
2422
  "eval_overall_precision": 0.9484536082474226,
2423
  "eval_overall_recall": 0.9387755102040817,
2424
+ "eval_runtime": 0.2951,
2425
+ "eval_samples_per_second": 576.156,
2426
+ "eval_steps_per_second": 10.167,
2427
  "step": 7488
2428
  },
2429
  {
 
2452
  "eval_overall_f1": 0.9389312977099236,
2453
  "eval_overall_precision": 0.9365482233502538,
2454
  "eval_overall_recall": 0.9413265306122449,
2455
+ "eval_runtime": 0.2963,
2456
+ "eval_samples_per_second": 573.823,
2457
+ "eval_steps_per_second": 10.126,
2458
  "step": 7584
2459
  },
2460
  {
 
2483
  "eval_overall_f1": 0.9453621346886911,
2484
  "eval_overall_precision": 0.9417721518987342,
2485
  "eval_overall_recall": 0.9489795918367347,
2486
+ "eval_runtime": 0.2956,
2487
+ "eval_samples_per_second": 575.01,
2488
+ "eval_steps_per_second": 10.147,
2489
  "step": 7680
2490
  },
2491
  {
 
2514
  "eval_overall_f1": 0.937579617834395,
2515
  "eval_overall_precision": 0.9363867684478372,
2516
  "eval_overall_recall": 0.9387755102040817,
2517
+ "eval_runtime": 0.2941,
2518
+ "eval_samples_per_second": 577.956,
2519
+ "eval_steps_per_second": 10.199,
2520
  "step": 7776
2521
  },
2522
  {
 
2545
  "eval_overall_f1": 0.9413265306122449,
2546
  "eval_overall_precision": 0.9413265306122449,
2547
  "eval_overall_recall": 0.9413265306122449,
2548
+ "eval_runtime": 0.2975,
2549
+ "eval_samples_per_second": 571.366,
2550
+ "eval_steps_per_second": 10.083,
2551
  "step": 7872
2552
  },
2553
  {
 
2576
  "eval_overall_f1": 0.940127388535032,
2577
  "eval_overall_precision": 0.9389312977099237,
2578
  "eval_overall_recall": 0.9413265306122449,
2579
+ "eval_runtime": 0.2941,
2580
+ "eval_samples_per_second": 578.062,
2581
+ "eval_steps_per_second": 10.201,
2582
  "step": 7968
2583
  },
2584
  {
 
2607
  "eval_overall_f1": 0.9390862944162437,
2608
  "eval_overall_precision": 0.9343434343434344,
2609
  "eval_overall_recall": 0.9438775510204082,
2610
+ "eval_runtime": 0.2949,
2611
+ "eval_samples_per_second": 576.398,
2612
+ "eval_steps_per_second": 10.172,
2613
  "step": 8064
2614
  },
2615
  {
 
2638
  "eval_overall_f1": 0.9441624365482234,
2639
  "eval_overall_precision": 0.9393939393939394,
2640
  "eval_overall_recall": 0.9489795918367347,
2641
+ "eval_runtime": 0.297,
2642
+ "eval_samples_per_second": 572.382,
2643
+ "eval_steps_per_second": 10.101,
2644
  "step": 8160
2645
  },
2646
  {
 
2669
  "eval_overall_f1": 0.9402795425667091,
2670
  "eval_overall_precision": 0.9367088607594937,
2671
  "eval_overall_recall": 0.9438775510204082,
2672
+ "eval_runtime": 0.2956,
2673
+ "eval_samples_per_second": 575.107,
2674
+ "eval_steps_per_second": 10.149,
2675
  "step": 8256
2676
  },
2677
  {
 
2700
  "eval_overall_f1": 0.9377382465057178,
2701
  "eval_overall_precision": 0.9341772151898734,
2702
  "eval_overall_recall": 0.9413265306122449,
2703
+ "eval_runtime": 0.2986,
2704
+ "eval_samples_per_second": 569.242,
2705
+ "eval_steps_per_second": 10.045,
2706
  "step": 8352
2707
  },
2708
  {
 
2731
  "eval_overall_f1": 0.935361216730038,
2732
  "eval_overall_precision": 0.929471032745592,
2733
  "eval_overall_recall": 0.9413265306122449,
2734
+ "eval_runtime": 0.2991,
2735
+ "eval_samples_per_second": 568.332,
2736
+ "eval_steps_per_second": 10.029,
2737
  "step": 8448
2738
  },
2739
  {
 
2762
  "eval_overall_f1": 0.9438775510204082,
2763
  "eval_overall_precision": 0.9438775510204082,
2764
  "eval_overall_recall": 0.9438775510204082,
2765
+ "eval_runtime": 0.3032,
2766
+ "eval_samples_per_second": 560.764,
2767
+ "eval_steps_per_second": 9.896,
2768
  "step": 8544
2769
  },
2770
  {
 
2793
  "eval_overall_f1": 0.9314720812182741,
2794
  "eval_overall_precision": 0.9267676767676768,
2795
  "eval_overall_recall": 0.9362244897959183,
2796
+ "eval_runtime": 0.2936,
2797
+ "eval_samples_per_second": 579.036,
2798
+ "eval_steps_per_second": 10.218,
2799
  "step": 8640
2800
  },
2801
  {
 
2824
  "eval_overall_f1": 0.9404309252217997,
2825
  "eval_overall_precision": 0.9345088161209067,
2826
  "eval_overall_recall": 0.9464285714285714,
2827
+ "eval_runtime": 0.2966,
2828
+ "eval_samples_per_second": 573.166,
2829
+ "eval_steps_per_second": 10.115,
2830
  "step": 8736
2831
  },
2832
  {
 
2855
  "eval_overall_f1": 0.935361216730038,
2856
  "eval_overall_precision": 0.929471032745592,
2857
  "eval_overall_recall": 0.9413265306122449,
2858
+ "eval_runtime": 0.2952,
2859
+ "eval_samples_per_second": 575.915,
2860
+ "eval_steps_per_second": 10.163,
2861
  "step": 8832
2862
  },
2863
  {
 
2886
  "eval_overall_f1": 0.935361216730038,
2887
  "eval_overall_precision": 0.929471032745592,
2888
  "eval_overall_recall": 0.9413265306122449,
2889
+ "eval_runtime": 0.2966,
2890
+ "eval_samples_per_second": 573.134,
2891
+ "eval_steps_per_second": 10.114,
2892
  "step": 8928
2893
  },
2894
  {
 
2917
  "eval_overall_f1": 0.9416243654822335,
2918
  "eval_overall_precision": 0.9368686868686869,
2919
  "eval_overall_recall": 0.9464285714285714,
2920
+ "eval_runtime": 0.2948,
2921
+ "eval_samples_per_second": 576.719,
2922
+ "eval_steps_per_second": 10.177,
2923
  "step": 9024
2924
  },
2925
  {
 
2948
  "eval_overall_f1": 0.9378960709759189,
2949
  "eval_overall_precision": 0.9319899244332494,
2950
  "eval_overall_recall": 0.9438775510204082,
2951
+ "eval_runtime": 0.295,
2952
+ "eval_samples_per_second": 576.364,
2953
+ "eval_steps_per_second": 10.171,
2954
  "step": 9120
2955
  },
2956
  {
 
2979
  "eval_overall_f1": 0.9341772151898734,
2980
  "eval_overall_precision": 0.9271356783919598,
2981
  "eval_overall_recall": 0.9413265306122449,
2982
+ "eval_runtime": 0.2942,
2983
+ "eval_samples_per_second": 577.85,
2984
+ "eval_steps_per_second": 10.197,
2985
  "step": 9216
2986
  },
2987
  {
 
3010
  "eval_overall_f1": 0.9416243654822335,
3011
  "eval_overall_precision": 0.9368686868686869,
3012
  "eval_overall_recall": 0.9464285714285714,
3013
+ "eval_runtime": 0.2974,
3014
+ "eval_samples_per_second": 571.562,
3015
+ "eval_steps_per_second": 10.086,
3016
  "step": 9312
3017
  },
3018
  {
 
3041
  "eval_overall_f1": 0.9390862944162437,
3042
  "eval_overall_precision": 0.9343434343434344,
3043
  "eval_overall_recall": 0.9438775510204082,
3044
+ "eval_runtime": 0.2935,
3045
+ "eval_samples_per_second": 579.188,
3046
+ "eval_steps_per_second": 10.221,
3047
  "step": 9408
3048
  },
3049
  {
 
3072
  "eval_overall_f1": 0.9390862944162437,
3073
  "eval_overall_precision": 0.9343434343434344,
3074
  "eval_overall_recall": 0.9438775510204082,
3075
+ "eval_runtime": 0.2978,
3076
+ "eval_samples_per_second": 570.864,
3077
+ "eval_steps_per_second": 10.074,
3078
  "step": 9504
3079
  },
3080
  {
 
3103
  "eval_overall_f1": 0.9390862944162437,
3104
  "eval_overall_precision": 0.9343434343434344,
3105
  "eval_overall_recall": 0.9438775510204082,
3106
+ "eval_runtime": 0.2953,
3107
+ "eval_samples_per_second": 575.755,
3108
+ "eval_steps_per_second": 10.16,
3109
  "step": 9600
3110
  },
3111
  {
 
3113
  "step": 9600,
3114
  "total_flos": 4315798421360676.0,
3115
  "train_loss": 0.03753832100580136,
3116
+ "train_runtime": 578.6389,
3117
+ "train_samples_per_second": 264.586,
3118
+ "train_steps_per_second": 16.591
3119
  }
3120
  ],
3121
  "logging_steps": 500,