File size: 7,873 Bytes
15cdfb2
5f432f5
15cdfb2
 
b2fe4a9
82c5d46
15cdfb2
ae98286
b9f5fd0
b2fe4a9
b652a35
 
06b312c
a3d2cd6
9e455e8
82c5d46
b2fe4a9
06b312c
b2fe4a9
ae98286
b9f5fd0
b2fe4a9
b652a35
b2fe4a9
 
 
 
8cf837b
dca759b
5f432f5
b2fe4a9
ae98286
b2fe4a9
 
82c5d46
b2fe4a9
b9f5fd0
b652a35
73f8345
b652a35
15cdfb2
b2fe4a9
 
8cf837b
73f8345
5f432f5
b2fe4a9
 
 
 
5f432f5
06b312c
9e455e8
4a31726
ae98286
26d96ac
 
ae98286
b2fe4a9
ae98286
5f432f5
b652a35
15cdfb2
b652a35
ae98286
b2fe4a9
4fe2063
06b312c
794e32b
b2fe4a9
 
4fe2063
9e455e8
4fe2063
b652a35
b2fe4a9
 
ae98286
b2fe4a9
 
 
26d96ac
b2fe4a9
b652a35
15cdfb2
b2fe4a9
 
 
a3d2cd6
15cdfb2
9e455e8
dca759b
b9f5fd0
b2fe4a9
ef4bf75
5f432f5
5b3238e
b2fe4a9
9e455e8
b2fe4a9
06b312c
ae98286
9e455e8
15cdfb2
bd7566e
b2fe4a9
15cdfb2
06b312c
 
ef4bf75
 
4fe2063
b2fe4a9
ae98286
b2fe4a9
 
 
8cf837b
5f432f5
ae98286
 
15cdfb2
06b312c
4fe2063
15cdfb2
ae98286
dca759b
5f432f5
b2fe4a9
dca759b
b652a35
4fe2063
 
ae98286
 
b2fe4a9
9e455e8
b652a35
4fe2063
15cdfb2
44ba8ed
9e455e8
5b3238e
15cdfb2
06b312c
15cdfb2
b652a35
 
b2fe4a9
73f8345
b2fe4a9
dca759b
82c5d46
9e455e8
 
 
 
ae98286
b2fe4a9
 
9e455e8
b9f5fd0
b2fe4a9
 
dca759b
5f432f5
b2fe4a9
 
 
dca759b
b652a35
15cdfb2
dca759b
b652a35
4fe2063
b2fe4a9
794e32b
5f432f5
b652a35
ae98286
 
a3d2cd6
794e32b
b2fe4a9
bd7566e
ae98286
15cdfb2
06b312c
b2fe4a9
ae98286
5b3238e
b2fe4a9
 
 
 
 
 
06b312c
b2fe4a9
82c5d46
5f432f5
b2fe4a9
 
5f432f5
b2fe4a9
82c5d46
b652a35
b2fe4a9
06b312c
b2fe4a9
4fe2063
ae98286
b2fe4a9
82c5d46
ae98286
b2fe4a9
 
ae98286
b2fe4a9
 
b652a35
06b312c
9e455e8
a3d2cd6
ae98286
4fe2063
ae98286
b652a35
8cf837b
b652a35
26d96ac
06b312c
ae98286
794e32b
9e455e8
dca759b
bd7566e
794e32b
ae98286
b2fe4a9
ae98286
b2fe4a9
 
 
 
5f432f5
b2fe4a9
44ba8ed
5f432f5
 
9e455e8
b2fe4a9
b652a35
06b312c
a3d2cd6
9e455e8
4fe2063
ae98286
5b3238e
15cdfb2
 
 
 
 
 
 
344c140
 
15cdfb2
 
 
 
 
 
344c140
15cdfb2
b2fe4a9
15cdfb2
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
94
95
96
97
98
99
100
101
102
103
104
105
106
107
108
109
110
111
112
113
114
115
116
117
118
119
120
121
122
123
124
125
126
127
128
129
130
131
132
133
134
135
136
137
138
139
140
141
142
143
144
145
146
147
148
149
150
151
152
153
154
155
156
157
158
159
160
161
162
163
164
165
166
167
168
169
170
171
172
173
174
175
176
177
178
179
180
181
182
183
184
185
186
187
188
189
190
191
192
193
194
195
196
197
198
199
200
201
202
203
204
205
206
207
208
209
210
211
212
213
214
215
216
217
218
219
220
221
222
223
224
225
226
227
228
229
230
231
232
233
234
235
236
237
238
239
240
241
242
243
244
245
246
247
248
249
250
251
252
253
254
255
256
257
258
259
260
261
262
263
264
265
266
267
268
269
270
271
272
273
274
275
276
277
278
279
280
281
282
283
284
285
286
287
288
289
290
291
292
293
294
{
  "_name_or_path": "distributed/llama-1b",
  "all_reduce_scores": {
    "0": "NOT_ALIVE",
    "1": "NON_PARTICIPATING",
    "10": "NON_PARTICIPATING",
    "100": "NOT_ALIVE",
    "101": "NON_PARTICIPATING",
    "102": "NON_PARTICIPATING",
    "103": "NON_PARTICIPATING",
    "104": "NON_PARTICIPATING",
    "105": "NON_PARTICIPATING",
    "106": "NON_PARTICIPATING",
    "107": "NOT_ALIVE",
    "108": "NOT_ALIVE",
    "109": "NOT_ALIVE",
    "11": "FAIL",
    "110": "NON_PARTICIPATING",
    "111": "NON_PARTICIPATING",
    "112": "FAIL",
    "113": "NOT_ALIVE",
    "114": "NON_PARTICIPATING",
    "115": "NON_PARTICIPATING",
    "116": "NOT_ALIVE",
    "117": "FAIL",
    "118": "NOT_ALIVE",
    "119": "NOT_ALIVE",
    "12": "NON_PARTICIPATING",
    "120": "NON_PARTICIPATING",
    "121": "NOT_ALIVE",
    "122": "NON_PARTICIPATING",
    "123": "NOT_ALIVE",
    "124": "NOT_ALIVE",
    "125": "NOT_ALIVE",
    "126": "NOT_ALIVE",
    "127": "NOT_ALIVE",
    "128": "NOT_ALIVE",
    "129": "NON_PARTICIPATING",
    "13": "NON_PARTICIPATING",
    "130": "NON_PARTICIPATING",
    "131": "NOT_ALIVE",
    "132": "NON_PARTICIPATING",
    "133": "FAIL",
    "134": "NON_PARTICIPATING",
    "135": "NON_PARTICIPATING",
    "136": "NON_PARTICIPATING",
    "137": "SUCCESS",
    "138": "FAIL",
    "139": "NON_PARTICIPATING",
    "14": "SUCCESS",
    "140": "NOT_ALIVE",
    "141": "NOT_ALIVE",
    "142": "NON_PARTICIPATING",
    "143": "NON_PARTICIPATING",
    "144": "NON_PARTICIPATING",
    "145": "NON_PARTICIPATING",
    "146": "NON_PARTICIPATING",
    "147": "NON_PARTICIPATING",
    "148": "NON_PARTICIPATING",
    "149": "NOT_ALIVE",
    "15": "NON_PARTICIPATING",
    "150": "NON_PARTICIPATING",
    "151": "NOT_ALIVE",
    "152": "NOT_ALIVE",
    "153": "FAIL",
    "154": "NON_PARTICIPATING",
    "155": "NON_PARTICIPATING",
    "156": "NON_PARTICIPATING",
    "157": "NON_PARTICIPATING",
    "158": "NOT_ALIVE",
    "159": "NOT_ALIVE",
    "16": "NOT_ALIVE",
    "160": "NON_PARTICIPATING",
    "161": "NON_PARTICIPATING",
    "162": "NON_PARTICIPATING",
    "163": "NOT_ALIVE",
    "164": "FAIL",
    "165": "FAIL",
    "166": "NON_PARTICIPATING",
    "167": "NOT_ALIVE",
    "168": "NON_PARTICIPATING",
    "169": "NOT_ALIVE",
    "17": "NOT_ALIVE",
    "170": "NON_PARTICIPATING",
    "171": "NOT_ALIVE",
    "172": "NON_PARTICIPATING",
    "173": "NON_PARTICIPATING",
    "174": "NOT_ALIVE",
    "175": "NOT_ALIVE",
    "176": "NOT_ALIVE",
    "177": "NOT_ALIVE",
    "178": "NON_PARTICIPATING",
    "179": "NON_PARTICIPATING",
    "18": "NOT_ALIVE",
    "180": "NOT_ALIVE",
    "181": "NOT_ALIVE",
    "182": "NOT_ALIVE",
    "183": "SUCCESS",
    "184": "NON_PARTICIPATING",
    "185": "NON_PARTICIPATING",
    "186": "NON_PARTICIPATING",
    "187": "NON_PARTICIPATING",
    "188": "NOT_ALIVE",
    "189": "NOT_ALIVE",
    "19": "NOT_ALIVE",
    "190": "NOT_ALIVE",
    "191": "NOT_ALIVE",
    "192": "NON_PARTICIPATING",
    "193": "NON_PARTICIPATING",
    "194": "NON_PARTICIPATING",
    "195": "NON_PARTICIPATING",
    "196": "NON_PARTICIPATING",
    "197": "NON_PARTICIPATING",
    "198": "NON_PARTICIPATING",
    "199": "NON_PARTICIPATING",
    "2": "NOT_ALIVE",
    "20": "NON_PARTICIPATING",
    "200": "NOT_ALIVE",
    "201": "NOT_ALIVE",
    "202": "NOT_ALIVE",
    "203": "NOT_ALIVE",
    "204": "NON_PARTICIPATING",
    "205": "NON_PARTICIPATING",
    "206": "NON_PARTICIPATING",
    "207": "NOT_ALIVE",
    "208": "NON_PARTICIPATING",
    "209": "NOT_ALIVE",
    "21": "NOT_ALIVE",
    "210": "NON_PARTICIPATING",
    "211": "NON_PARTICIPATING",
    "212": "NON_PARTICIPATING",
    "213": "NON_PARTICIPATING",
    "214": "NON_PARTICIPATING",
    "215": "FAIL",
    "216": "NON_PARTICIPATING",
    "217": "NOT_ALIVE",
    "218": "NON_PARTICIPATING",
    "219": "NON_PARTICIPATING",
    "22": "NON_PARTICIPATING",
    "220": "NON_PARTICIPATING",
    "221": "NOT_ALIVE",
    "222": "NOT_ALIVE",
    "223": "NON_PARTICIPATING",
    "224": "NOT_ALIVE",
    "225": "NON_PARTICIPATING",
    "226": "NOT_ALIVE",
    "227": "NON_PARTICIPATING",
    "228": "NON_PARTICIPATING",
    "229": "NON_PARTICIPATING",
    "23": "NON_PARTICIPATING",
    "230": "NOT_ALIVE",
    "231": "NOT_ALIVE",
    "232": "NOT_ALIVE",
    "233": "NON_PARTICIPATING",
    "234": "NON_PARTICIPATING",
    "235": "NON_PARTICIPATING",
    "236": "NOT_ALIVE",
    "237": "NOT_ALIVE",
    "238": "SUCCESS",
    "239": "NOT_ALIVE",
    "24": "NON_PARTICIPATING",
    "240": "NOT_ALIVE",
    "241": "NON_PARTICIPATING",
    "242": "NOT_ALIVE",
    "243": "NOT_ALIVE",
    "244": "NON_PARTICIPATING",
    "245": "NON_PARTICIPATING",
    "246": "NOT_ALIVE",
    "247": "NOT_ALIVE",
    "248": "NOT_ALIVE",
    "249": "NOT_ALIVE",
    "25": "SUCCESS",
    "250": "NON_PARTICIPATING",
    "251": "NOT_ALIVE",
    "252": "NON_PARTICIPATING",
    "253": "NOT_ALIVE",
    "254": "NON_PARTICIPATING",
    "255": "NOT_ALIVE",
    "26": "NON_PARTICIPATING",
    "27": "NOT_ALIVE",
    "28": "NON_PARTICIPATING",
    "29": "NON_PARTICIPATING",
    "3": "NON_PARTICIPATING",
    "30": "NOT_ALIVE",
    "31": "NON_PARTICIPATING",
    "32": "FAIL",
    "33": "NOT_ALIVE",
    "34": "NOT_ALIVE",
    "35": "NON_PARTICIPATING",
    "36": "NON_PARTICIPATING",
    "37": "NOT_ALIVE",
    "38": "SUCCESS",
    "39": "FAIL",
    "4": "NOT_ALIVE",
    "40": "NON_PARTICIPATING",
    "41": "SUCCESS",
    "42": "NON_PARTICIPATING",
    "43": "NON_PARTICIPATING",
    "44": "NON_PARTICIPATING",
    "45": "NOT_ALIVE",
    "46": "NOT_ALIVE",
    "47": "NOT_ALIVE",
    "48": "NOT_ALIVE",
    "49": "NON_PARTICIPATING",
    "5": "NON_PARTICIPATING",
    "50": "NON_PARTICIPATING",
    "51": "NON_PARTICIPATING",
    "52": "NON_PARTICIPATING",
    "53": "NON_PARTICIPATING",
    "54": "NOT_ALIVE",
    "55": "NOT_ALIVE",
    "56": "FAIL",
    "57": "NON_PARTICIPATING",
    "58": "NON_PARTICIPATING",
    "59": "NON_PARTICIPATING",
    "6": "NOT_ALIVE",
    "60": "NOT_ALIVE",
    "61": "NON_PARTICIPATING",
    "62": "NON_PARTICIPATING",
    "63": "FAIL",
    "64": "NON_PARTICIPATING",
    "65": "NON_PARTICIPATING",
    "66": "NON_PARTICIPATING",
    "67": "NOT_ALIVE",
    "68": "SUCCESS",
    "69": "NON_PARTICIPATING",
    "7": "NOT_ALIVE",
    "70": "NON_PARTICIPATING",
    "71": "NON_PARTICIPATING",
    "72": "NON_PARTICIPATING",
    "73": "NON_PARTICIPATING",
    "74": "NOT_ALIVE",
    "75": "NON_PARTICIPATING",
    "76": "NON_PARTICIPATING",
    "77": "NON_PARTICIPATING",
    "78": "NON_PARTICIPATING",
    "79": "NON_PARTICIPATING",
    "8": "NON_PARTICIPATING",
    "80": "NON_PARTICIPATING",
    "81": "NON_PARTICIPATING",
    "82": "NON_PARTICIPATING",
    "83": "NON_PARTICIPATING",
    "84": "NON_PARTICIPATING",
    "85": "NOT_ALIVE",
    "86": "NOT_ALIVE",
    "87": "NON_PARTICIPATING",
    "88": "NON_PARTICIPATING",
    "89": "NOT_ALIVE",
    "9": "NOT_ALIVE",
    "90": "NON_PARTICIPATING",
    "91": "NOT_ALIVE",
    "92": "NOT_ALIVE",
    "93": "NOT_ALIVE",
    "94": "NON_PARTICIPATING",
    "95": "NOT_ALIVE",
    "96": "NOT_ALIVE",
    "97": "NON_PARTICIPATING",
    "98": "NON_PARTICIPATING",
    "99": "NOT_ALIVE"
  },
  "architectures": [
    "LlamaForCausalLM"
  ],
  "attention_bias": false,
  "attention_dropout": 0.0,
  "block_list": [
    5953074,
    5953089
  ],
  "bos_token_id": 1,
  "eos_token_id": 2,
  "hidden_act": "silu",
  "hidden_size": 2048,
  "initializer_range": 0.02,
  "inner_step": 41,
  "intermediate_size": 5632,
  "last_allreduce_block": 5951984,
  "max_position_embeddings": 2048,
  "mlp_bias": false,
  "model_type": "llama",
  "num_attention_heads": 32,
  "num_hidden_layers": 22,
  "num_key_value_heads": 4,
  "pretraining_tp": 1,
  "rms_norm_eps": 1e-05,
  "rope_scaling": null,
  "rope_theta": 10000.0,
  "tie_word_embeddings": false,
  "torch_dtype": "float32",
  "transformers_version": "4.39.3",
  "use_cache": false,
  "vocab_size": 32000
}