Kumarverma11 commited on
Commit
3aac170
·
verified ·
1 Parent(s): 7f409e0

Phase3 ckpt step 50

Browse files
checkpoints/checkpoint-50/adapter_model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:2c05602edd46e4cf53d8b2e2c7d9185cb9bb2dc759775117a05cd5c1589b2b27
3
  size 119801528
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:5f91cce9515723b65c6fafaacbdf9b3de96a50d42bbdebe0c8c7f30d1690f608
3
  size 119801528
checkpoints/checkpoint-50/optimizer.pt CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:33c489ad20e393d48e76a172c78aafc073b3482e6f5405fe9250551604367237
3
  size 61397701
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:14ea8b43b3ff7d5fd2f18f25cd83613ed615bee210ff94290823a8bd573ab146
3
  size 61397701
checkpoints/checkpoint-50/trainer_state.json CHANGED
@@ -11,44 +11,44 @@
11
  "log_history": [
12
  {
13
  "epoch": 0.019417475728155338,
14
- "grad_norm": 0.5592218637466431,
15
  "learning_rate": 0.0,
16
- "loss": 1.0037728548049927,
17
  "step": 1
18
  },
19
  {
20
  "epoch": 0.1941747572815534,
21
- "grad_norm": 0.6074398159980774,
22
  "learning_rate": 9.524135262330098e-05,
23
- "loss": 0.9681400722927518,
24
  "step": 10
25
  },
26
  {
27
  "epoch": 0.3883495145631068,
28
- "grad_norm": 0.6838647723197937,
29
  "learning_rate": 7.408768370508576e-05,
30
- "loss": 0.9677780151367188,
31
  "step": 20
32
  },
33
  {
34
  "epoch": 0.5825242718446602,
35
- "grad_norm": 0.6627008318901062,
36
  "learning_rate": 4.373333832178478e-05,
37
- "loss": 0.9933381080627441,
38
  "step": 30
39
  },
40
  {
41
  "epoch": 0.7766990291262136,
42
- "grad_norm": 0.6596543192863464,
43
  "learning_rate": 1.5772644703565565e-05,
44
- "loss": 0.9614946365356445,
45
  "step": 40
46
  },
47
  {
48
  "epoch": 0.970873786407767,
49
- "grad_norm": 0.6468517780303955,
50
  "learning_rate": 8.856374635655695e-07,
51
- "loss": 0.9476834297180176,
52
  "step": 50
53
  }
54
  ],
@@ -69,7 +69,7 @@
69
  "attributes": {}
70
  }
71
  },
72
- "total_flos": 6663849985376256.0,
73
  "train_batch_size": 8,
74
  "trial_name": null,
75
  "trial_params": null
 
11
  "log_history": [
12
  {
13
  "epoch": 0.019417475728155338,
14
+ "grad_norm": 0.6309196949005127,
15
  "learning_rate": 0.0,
16
+ "loss": 1.007138967514038,
17
  "step": 1
18
  },
19
  {
20
  "epoch": 0.1941747572815534,
21
+ "grad_norm": 0.6809913516044617,
22
  "learning_rate": 9.524135262330098e-05,
23
+ "loss": 0.9641060299343533,
24
  "step": 10
25
  },
26
  {
27
  "epoch": 0.3883495145631068,
28
+ "grad_norm": 0.6508110761642456,
29
  "learning_rate": 7.408768370508576e-05,
30
+ "loss": 0.9679458618164063,
31
  "step": 20
32
  },
33
  {
34
  "epoch": 0.5825242718446602,
35
+ "grad_norm": 0.6529077887535095,
36
  "learning_rate": 4.373333832178478e-05,
37
+ "loss": 0.956454086303711,
38
  "step": 30
39
  },
40
  {
41
  "epoch": 0.7766990291262136,
42
+ "grad_norm": 0.6825394630432129,
43
  "learning_rate": 1.5772644703565565e-05,
44
+ "loss": 0.9437665939331055,
45
  "step": 40
46
  },
47
  {
48
  "epoch": 0.970873786407767,
49
+ "grad_norm": 0.6889418959617615,
50
  "learning_rate": 8.856374635655695e-07,
51
+ "loss": 0.9414969444274902,
52
  "step": 50
53
  }
54
  ],
 
69
  "attributes": {}
70
  }
71
  },
72
+ "total_flos": 6649445008539648.0,
73
  "train_batch_size": 8,
74
  "trial_name": null,
75
  "trial_params": null
checkpoints/checkpoint-50/training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:f6ee30206dd9557cdb72e96f4e6d60f78dc1179da2c42c668f6f07ff70838e45
3
  size 5777
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:d7eb50cc13c649388d04f9d19ca831e6db6e0f12509fe17a7c1bb09fd8072213
3
  size 5777