hanngao commited on
Commit
150bb1d
·
verified ·
1 Parent(s): 423ed60

Training in progress, epoch 3, checkpoint

Browse files
last-checkpoint/model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:8361afeef393c81a02086ba51c86073ac81e86c12bc7610d2c6bdda710514d48
3
  size 346293856
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:a62f73f0d7e2123e0401bf70219629f37dba53bf4ea79d6e53d3e33bf079f52d
3
  size 346293856
last-checkpoint/trainer_state.json CHANGED
@@ -2,9 +2,9 @@
2
  "best_global_step": 16,
3
  "best_metric": 0.8921805620193481,
4
  "best_model_checkpoint": "./finetuning/checkpoint-16",
5
- "epoch": 2.0,
6
  "eval_steps": 500,
7
- "global_step": 32,
8
  "is_hyper_param_search": false,
9
  "is_local_process_zero": true,
10
  "is_world_process_zero": true,
@@ -47,6 +47,22 @@
47
  "eval_samples_per_second": 19.587,
48
  "eval_steps_per_second": 1.224,
49
  "step": 32
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
50
  }
51
  ],
52
  "logging_steps": 10,
@@ -61,7 +77,7 @@
61
  "early_stopping_threshold": 0.0
62
  },
63
  "attributes": {
64
- "early_stopping_patience_counter": 1
65
  }
66
  },
67
  "TrainerControl": {
@@ -70,12 +86,12 @@
70
  "should_evaluate": false,
71
  "should_log": false,
72
  "should_save": true,
73
- "should_training_stop": false
74
  },
75
  "attributes": {}
76
  }
77
  },
78
- "total_flos": 1.250962187747328e+18,
79
  "train_batch_size": 512,
80
  "trial_name": null,
81
  "trial_params": null
 
2
  "best_global_step": 16,
3
  "best_metric": 0.8921805620193481,
4
  "best_model_checkpoint": "./finetuning/checkpoint-16",
5
+ "epoch": 3.0,
6
  "eval_steps": 500,
7
+ "global_step": 48,
8
  "is_hyper_param_search": false,
9
  "is_local_process_zero": true,
10
  "is_world_process_zero": true,
 
47
  "eval_samples_per_second": 19.587,
48
  "eval_steps_per_second": 1.224,
49
  "step": 32
50
+ },
51
+ {
52
+ "epoch": 2.5,
53
+ "grad_norm": 0.8204277753829956,
54
+ "learning_rate": 3.78125e-05,
55
+ "loss": 0.2904698610305786,
56
+ "step": 40
57
+ },
58
+ {
59
+ "epoch": 3.0,
60
+ "eval_accuracy": 0.774,
61
+ "eval_loss": 0.9198980927467346,
62
+ "eval_runtime": 102.9304,
63
+ "eval_samples_per_second": 19.431,
64
+ "eval_steps_per_second": 1.214,
65
+ "step": 48
66
  }
67
  ],
68
  "logging_steps": 10,
 
77
  "early_stopping_threshold": 0.0
78
  },
79
  "attributes": {
80
+ "early_stopping_patience_counter": 2
81
  }
82
  },
83
  "TrainerControl": {
 
86
  "should_evaluate": false,
87
  "should_log": false,
88
  "should_save": true,
89
+ "should_training_stop": true
90
  },
91
  "attributes": {}
92
  }
93
  },
94
+ "total_flos": 1.876443281620992e+18,
95
  "train_batch_size": 512,
96
  "trial_name": null,
97
  "trial_params": null