crypt0trading commited on
Commit
a97fe6d
·
verified ·
1 Parent(s): b3ec351

Run 2. Outer Step 11. Inner Step 0.

Browse files
Files changed (3) hide show
  1. config.json +25 -31
  2. inner_optimizer.pt +1 -1
  3. model.safetensors +1 -1
config.json CHANGED
@@ -1,20 +1,20 @@
1
  {
2
- "_name_or_path": "crypt0trading/c66-h8",
3
  "activation_function": "gelu_new",
4
  "all_reduce_scores": {
5
  "0": "NON_PARTICIPATING",
6
- "1": "SUCCESS",
7
  "10": "NON_PARTICIPATING",
8
  "100": "NON_PARTICIPATING",
9
- "101": "NON_PARTICIPATING",
10
- "102": "NON_PARTICIPATING",
11
  "103": "NON_PARTICIPATING",
12
  "104": "NON_PARTICIPATING",
13
  "105": "NON_PARTICIPATING",
14
  "106": "NON_PARTICIPATING",
15
  "107": "NON_PARTICIPATING",
16
  "108": "NON_PARTICIPATING",
17
- "109": "SUCCESS",
18
  "11": "NON_PARTICIPATING",
19
  "110": "NON_PARTICIPATING",
20
  "111": "NON_PARTICIPATING",
@@ -23,7 +23,7 @@
23
  "114": "NON_PARTICIPATING",
24
  "115": "NON_PARTICIPATING",
25
  "116": "NON_PARTICIPATING",
26
- "117": "NON_PARTICIPATING",
27
  "118": "NON_PARTICIPATING",
28
  "119": "NON_PARTICIPATING",
29
  "12": "NON_PARTICIPATING",
@@ -75,7 +75,7 @@
75
  "161": "NON_PARTICIPATING",
76
  "162": "NON_PARTICIPATING",
77
  "163": "NON_PARTICIPATING",
78
- "164": "NON_PARTICIPATING",
79
  "165": "NON_PARTICIPATING",
80
  "166": "NON_PARTICIPATING",
81
  "167": "NON_PARTICIPATING",
@@ -89,7 +89,7 @@
89
  "174": "NON_PARTICIPATING",
90
  "175": "NON_PARTICIPATING",
91
  "176": "NON_PARTICIPATING",
92
- "177": "NON_PARTICIPATING",
93
  "178": "NON_PARTICIPATING",
94
  "179": "NON_PARTICIPATING",
95
  "18": "NON_PARTICIPATING",
@@ -106,8 +106,8 @@
106
  "19": "NON_PARTICIPATING",
107
  "190": "NON_PARTICIPATING",
108
  "191": "NON_PARTICIPATING",
109
- "192": "NON_PARTICIPATING",
110
- "193": "NON_PARTICIPATING",
111
  "194": "NON_PARTICIPATING",
112
  "195": "NON_PARTICIPATING",
113
  "196": "NON_PARTICIPATING",
@@ -139,12 +139,12 @@
139
  "219": "NON_PARTICIPATING",
140
  "22": "SUCCESS",
141
  "220": "NON_PARTICIPATING",
142
- "221": "NON_PARTICIPATING",
143
  "222": "NON_PARTICIPATING",
144
- "223": "NON_PARTICIPATING",
145
  "224": "NON_PARTICIPATING",
146
  "225": "NON_PARTICIPATING",
147
- "226": "NON_PARTICIPATING",
148
  "227": "NON_PARTICIPATING",
149
  "228": "NON_PARTICIPATING",
150
  "229": "NON_PARTICIPATING",
@@ -154,7 +154,7 @@
154
  "232": "NON_PARTICIPATING",
155
  "233": "NON_PARTICIPATING",
156
  "234": "NON_PARTICIPATING",
157
- "235": "NON_PARTICIPATING",
158
  "236": "NON_PARTICIPATING",
159
  "237": "NON_PARTICIPATING",
160
  "238": "NON_PARTICIPATING",
@@ -169,9 +169,9 @@
169
  "246": "NON_PARTICIPATING",
170
  "247": "NON_PARTICIPATING",
171
  "248": "NON_PARTICIPATING",
172
- "249": "NON_PARTICIPATING",
173
  "25": "SUCCESS",
174
- "250": "NON_PARTICIPATING",
175
  "251": "NON_PARTICIPATING",
176
  "252": "NON_PARTICIPATING",
177
  "253": "NON_PARTICIPATING",
@@ -191,12 +191,12 @@
191
  "36": "NON_PARTICIPATING",
192
  "37": "NON_PARTICIPATING",
193
  "38": "NON_PARTICIPATING",
194
- "39": "SUCCESS",
195
  "4": "NON_PARTICIPATING",
196
  "40": "NON_PARTICIPATING",
197
  "41": "NON_PARTICIPATING",
198
- "42": "NON_PARTICIPATING",
199
- "43": "NON_PARTICIPATING",
200
  "44": "NON_PARTICIPATING",
201
  "45": "NON_PARTICIPATING",
202
  "46": "NON_PARTICIPATING",
@@ -217,12 +217,12 @@
217
  "6": "NON_PARTICIPATING",
218
  "60": "NON_PARTICIPATING",
219
  "61": "NON_PARTICIPATING",
220
- "62": "SUCCESS",
221
  "63": "NON_PARTICIPATING",
222
  "64": "NON_PARTICIPATING",
223
  "65": "NON_PARTICIPATING",
224
  "66": "NON_PARTICIPATING",
225
- "67": "SUCCESS",
226
  "68": "NON_PARTICIPATING",
227
  "69": "NON_PARTICIPATING",
228
  "7": "NON_PARTICIPATING",
@@ -237,7 +237,7 @@
237
  "78": "NON_PARTICIPATING",
238
  "79": "NON_PARTICIPATING",
239
  "8": "NON_PARTICIPATING",
240
- "80": "NON_PARTICIPATING",
241
  "81": "NON_PARTICIPATING",
242
  "82": "NON_PARTICIPATING",
243
  "83": "NON_PARTICIPATING",
@@ -267,21 +267,15 @@
267
  "AutoConfig": "distributed/optimized-gpt2-500m--configuration_gpt_optimized.GPTOptimConfig",
268
  "AutoModelForCausalLM": "distributed/optimized-gpt2-500m--modeling_gpt_optimized.GPTOptim"
269
  },
270
- "block_list": [
271
- 5363652,
272
- 5363657,
273
- 5363661,
274
- 5363665,
275
- 5363669
276
- ],
277
  "block_size": 1024,
278
  "bos_token_id": 50256,
279
  "embd_pdrop": 0.1,
280
  "eos_token_id": 50256,
281
  "initializer_range": 0.02,
282
- "inner_step": 26,
283
  "inner_steps": 0,
284
- "last_allreduce_block": 5339896,
285
  "layer_norm_epsilon": 1e-05,
286
  "model_type": "gpt_optimized",
287
  "n_embd": 1280,
 
1
  {
2
+ "_name_or_path": "distributed/optimized-gpt2-1b",
3
  "activation_function": "gelu_new",
4
  "all_reduce_scores": {
5
  "0": "NON_PARTICIPATING",
6
+ "1": "NON_PARTICIPATING",
7
  "10": "NON_PARTICIPATING",
8
  "100": "NON_PARTICIPATING",
9
+ "101": "SUCCESS",
10
+ "102": "SUCCESS",
11
  "103": "NON_PARTICIPATING",
12
  "104": "NON_PARTICIPATING",
13
  "105": "NON_PARTICIPATING",
14
  "106": "NON_PARTICIPATING",
15
  "107": "NON_PARTICIPATING",
16
  "108": "NON_PARTICIPATING",
17
+ "109": "NON_PARTICIPATING",
18
  "11": "NON_PARTICIPATING",
19
  "110": "NON_PARTICIPATING",
20
  "111": "NON_PARTICIPATING",
 
23
  "114": "NON_PARTICIPATING",
24
  "115": "NON_PARTICIPATING",
25
  "116": "NON_PARTICIPATING",
26
+ "117": "SUCCESS",
27
  "118": "NON_PARTICIPATING",
28
  "119": "NON_PARTICIPATING",
29
  "12": "NON_PARTICIPATING",
 
75
  "161": "NON_PARTICIPATING",
76
  "162": "NON_PARTICIPATING",
77
  "163": "NON_PARTICIPATING",
78
+ "164": "SUCCESS",
79
  "165": "NON_PARTICIPATING",
80
  "166": "NON_PARTICIPATING",
81
  "167": "NON_PARTICIPATING",
 
89
  "174": "NON_PARTICIPATING",
90
  "175": "NON_PARTICIPATING",
91
  "176": "NON_PARTICIPATING",
92
+ "177": "SUCCESS",
93
  "178": "NON_PARTICIPATING",
94
  "179": "NON_PARTICIPATING",
95
  "18": "NON_PARTICIPATING",
 
106
  "19": "NON_PARTICIPATING",
107
  "190": "NON_PARTICIPATING",
108
  "191": "NON_PARTICIPATING",
109
+ "192": "SUCCESS",
110
+ "193": "SUCCESS",
111
  "194": "NON_PARTICIPATING",
112
  "195": "NON_PARTICIPATING",
113
  "196": "NON_PARTICIPATING",
 
139
  "219": "NON_PARTICIPATING",
140
  "22": "SUCCESS",
141
  "220": "NON_PARTICIPATING",
142
+ "221": "SUCCESS",
143
  "222": "NON_PARTICIPATING",
144
+ "223": "FAIL",
145
  "224": "NON_PARTICIPATING",
146
  "225": "NON_PARTICIPATING",
147
+ "226": "SUCCESS",
148
  "227": "NON_PARTICIPATING",
149
  "228": "NON_PARTICIPATING",
150
  "229": "NON_PARTICIPATING",
 
154
  "232": "NON_PARTICIPATING",
155
  "233": "NON_PARTICIPATING",
156
  "234": "NON_PARTICIPATING",
157
+ "235": "SUCCESS",
158
  "236": "NON_PARTICIPATING",
159
  "237": "NON_PARTICIPATING",
160
  "238": "NON_PARTICIPATING",
 
169
  "246": "NON_PARTICIPATING",
170
  "247": "NON_PARTICIPATING",
171
  "248": "NON_PARTICIPATING",
172
+ "249": "SUCCESS",
173
  "25": "SUCCESS",
174
+ "250": "SUCCESS",
175
  "251": "NON_PARTICIPATING",
176
  "252": "NON_PARTICIPATING",
177
  "253": "NON_PARTICIPATING",
 
191
  "36": "NON_PARTICIPATING",
192
  "37": "NON_PARTICIPATING",
193
  "38": "NON_PARTICIPATING",
194
+ "39": "NON_PARTICIPATING",
195
  "4": "NON_PARTICIPATING",
196
  "40": "NON_PARTICIPATING",
197
  "41": "NON_PARTICIPATING",
198
+ "42": "SUCCESS",
199
+ "43": "SUCCESS",
200
  "44": "NON_PARTICIPATING",
201
  "45": "NON_PARTICIPATING",
202
  "46": "NON_PARTICIPATING",
 
217
  "6": "NON_PARTICIPATING",
218
  "60": "NON_PARTICIPATING",
219
  "61": "NON_PARTICIPATING",
220
+ "62": "NON_PARTICIPATING",
221
  "63": "NON_PARTICIPATING",
222
  "64": "NON_PARTICIPATING",
223
  "65": "NON_PARTICIPATING",
224
  "66": "NON_PARTICIPATING",
225
+ "67": "NON_PARTICIPATING",
226
  "68": "NON_PARTICIPATING",
227
  "69": "NON_PARTICIPATING",
228
  "7": "NON_PARTICIPATING",
 
237
  "78": "NON_PARTICIPATING",
238
  "79": "NON_PARTICIPATING",
239
  "8": "NON_PARTICIPATING",
240
+ "80": "SUCCESS",
241
  "81": "NON_PARTICIPATING",
242
  "82": "NON_PARTICIPATING",
243
  "83": "NON_PARTICIPATING",
 
267
  "AutoConfig": "distributed/optimized-gpt2-500m--configuration_gpt_optimized.GPTOptimConfig",
268
  "AutoModelForCausalLM": "distributed/optimized-gpt2-500m--modeling_gpt_optimized.GPTOptim"
269
  },
270
+ "block_list": [],
 
 
 
 
 
 
271
  "block_size": 1024,
272
  "bos_token_id": 50256,
273
  "embd_pdrop": 0.1,
274
  "eos_token_id": 50256,
275
  "initializer_range": 0.02,
276
+ "inner_step": 0,
277
  "inner_steps": 0,
278
+ "last_allreduce_block": 5363508,
279
  "layer_norm_epsilon": 1e-05,
280
  "model_type": "gpt_optimized",
281
  "n_embd": 1280,
inner_optimizer.pt CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:2275be2a61fb29cfe39f8ef164b66708d490ae891a5db06fa5b21a07368f4f71
3
  size 8081782026
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e1e7676cf1cd9ec2c4a00c238d39c1f344ee5a7dd27ad65e72ab3c13419f95fc
3
  size 8081782026
model.safetensors CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:bd989a8c8b2c1a3caa8adbc358659ff57c66096a808a9e6582e8ce3bfbaa24b2
3
  size 4040701744
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:7ba59bf6f504d702d1349e92b0cf12064e903119ec0fd54b44f5b0348e647f9e
3
  size 4040701744