File size: 7,821 Bytes
fed1dec
 
 
 
b9b3d26
857ae7e
fed1dec
25666dd
70278ac
 
e9e9a6c
 
 
4ff875e
25666dd
58e7ca8
9c4b308
25666dd
e9e9a6c
25666dd
9203fdb
f4c5d41
e9e9a6c
70278ac
25666dd
389b27a
70278ac
b9b3d26
70278ac
0ed2b3d
25666dd
a5ed109
25666dd
351df8d
f4c5d41
389b27a
9203fdb
f4c5d41
1f2bc56
b9b3d26
fed1dec
73e26fe
25666dd
f4c5d41
25666dd
 
 
 
f4c5d41
351df8d
e9e9a6c
70278ac
ad97c67
351df8d
25666dd
 
f4c5d41
e9e9a6c
25666dd
33319d7
857ae7e
25666dd
fed1dec
1cc0367
f4c5d41
e9e9a6c
25666dd
1f2bc56
e9e9a6c
25666dd
351df8d
edfc2e8
f4c5d41
25666dd
 
389b27a
25666dd
 
 
 
 
87c3655
389b27a
ad3e11c
ee90468
70278ac
25666dd
e9e9a6c
4ff875e
fed1dec
7f6c7b7
36d3e06
857ae7e
389b27a
351df8d
0ed2b3d
1b33ab7
ad97c67
e9e9a6c
b9b3d26
351df8d
70278ac
ad97c67
fed1dec
5e24562
951ae4a
fed1dec
70278ac
 
e9e9a6c
 
f4c5d41
25666dd
70278ac
25666dd
70278ac
b9b3d26
626fd28
ad97c67
a5ed109
 
f4c5d41
25666dd
f4c5d41
fed1dec
e9e9a6c
70278ac
0ed2b3d
24efb87
e9e9a6c
25666dd
 
70278ac
25666dd
e9e9a6c
9c4b308
25666dd
70278ac
 
e9e9a6c
1b33ab7
7f6c7b7
b9b3d26
fed1dec
33319d7
4cf6732
70278ac
f4c5d41
e9e9a6c
951ae4a
25666dd
ad97c67
58e7ca8
73e26fe
70278ac
351df8d
7f6c7b7
f4c5d41
25666dd
389b27a
951ae4a
9203fdb
25666dd
70278ac
36d3e06
ad3e11c
25666dd
ad97c67
73e26fe
36d3e06
1cc0367
fed1dec
f4c5d41
1cc0367
351df8d
389b27a
25666dd
 
f4c5d41
70278ac
 
b9b3d26
70d431f
f4c5d41
951ae4a
25666dd
fed1dec
25666dd
 
fb58348
5e24562
ad97c67
25666dd
70278ac
f4c5d41
951ae4a
b9b3d26
25666dd
 
58e7ca8
0ed2b3d
25666dd
 
951ae4a
25666dd
70278ac
 
25666dd
e9e9a6c
25666dd
b40bd48
951ae4a
edfc2e8
25666dd
 
b9b3d26
951ae4a
e9e9a6c
70278ac
25666dd
70278ac
25666dd
857ae7e
4ff875e
951ae4a
70278ac
b9b3d26
e9e9a6c
 
70278ac
ad97c67
f4c5d41
25666dd
 
ad97c67
857ae7e
fb58348
857ae7e
25666dd
70278ac
73e26fe
70278ac
fb58348
e9e9a6c
951ae4a
ad97c67
f4c5d41
5e24562
351df8d
f4c5d41
25666dd
 
f4c5d41
25666dd
4ff875e
73e26fe
25666dd
6d0ac76
dcd974d
fed1dec
 
 
 
 
 
 
3f65027
 
fed1dec
 
 
 
 
 
3f65027
fed1dec
25666dd
fed1dec
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
40
41
42
43
44
45
46
47
48
49
50
51
52
53
54
55
56
57
58
59
60
61
62
63
64
65
66
67
68
69
70
71
72
73
74
75
76
77
78
79
80
81
82
83
84
85
86
87
88
89
90
91
92
93
94
95
96
97
98
99
100
101
102
103
104
105
106
107
108
109
110
111
112
113
114
115
116
117
118
119
120
121
122
123
124
125
126
127
128
129
130
131
132
133
134
135
136
137
138
139
140
141
142
143
144
145
146
147
148
149
150
151
152
153
154
155
156
157
158
159
160
161
162
163
164
165
166
167
168
169
170
171
172
173
174
175
176
177
178
179
180
181
182
183
184
185
186
187
188
189
190
191
192
193
194
195
196
197
198
199
200
201
202
203
204
205
206
207
208
209
210
211
212
213
214
215
216
217
218
219
220
221
222
223
224
225
226
227
228
229
230
231
232
233
234
235
236
237
238
239
240
241
242
243
244
245
246
247
248
249
250
251
252
253
254
255
256
257
258
259
260
261
262
263
264
265
266
267
268
269
270
271
272
273
274
275
276
277
278
279
280
281
282
283
284
285
286
287
288
289
290
291
292
293
294
{
  "_name_or_path": "distributed/llama-1b",
  "all_reduce_scores": {
    "0": "NOT_ALIVE",
    "1": "NOT_ALIVE",
    "10": "NOT_ALIVE",
    "100": "NOT_ALIVE",
    "101": "NON_PARTICIPATING",
    "102": "NOT_ALIVE",
    "103": "NON_PARTICIPATING",
    "104": "NON_PARTICIPATING",
    "105": "NON_PARTICIPATING",
    "106": "NON_PARTICIPATING",
    "107": "NOT_ALIVE",
    "108": "NON_PARTICIPATING",
    "109": "NOT_ALIVE",
    "11": "NOT_ALIVE",
    "110": "NOT_ALIVE",
    "111": "NON_PARTICIPATING",
    "112": "NON_PARTICIPATING",
    "113": "NOT_ALIVE",
    "114": "NON_PARTICIPATING",
    "115": "NON_PARTICIPATING",
    "116": "NON_PARTICIPATING",
    "117": "NON_PARTICIPATING",
    "118": "NOT_ALIVE",
    "119": "NOT_ALIVE",
    "12": "NON_PARTICIPATING",
    "120": "NOT_ALIVE",
    "121": "NOT_ALIVE",
    "122": "NOT_ALIVE",
    "123": "NOT_ALIVE",
    "124": "NON_PARTICIPATING",
    "125": "NON_PARTICIPATING",
    "126": "NON_PARTICIPATING",
    "127": "NOT_ALIVE",
    "128": "NOT_ALIVE",
    "129": "NON_PARTICIPATING",
    "13": "NOT_ALIVE",
    "130": "NOT_ALIVE",
    "131": "NOT_ALIVE",
    "132": "NON_PARTICIPATING",
    "133": "NON_PARTICIPATING",
    "134": "NON_PARTICIPATING",
    "135": "NOT_ALIVE",
    "136": "NON_PARTICIPATING",
    "137": "NON_PARTICIPATING",
    "138": "SUCCESS",
    "139": "NON_PARTICIPATING",
    "14": "NON_PARTICIPATING",
    "140": "NON_PARTICIPATING",
    "141": "NON_PARTICIPATING",
    "142": "NON_PARTICIPATING",
    "143": "NON_PARTICIPATING",
    "144": "NOT_ALIVE",
    "145": "NON_PARTICIPATING",
    "146": "NON_PARTICIPATING",
    "147": "NON_PARTICIPATING",
    "148": "SUCCESS",
    "149": "NON_PARTICIPATING",
    "15": "NOT_ALIVE",
    "150": "NON_PARTICIPATING",
    "151": "NOT_ALIVE",
    "152": "NOT_ALIVE",
    "153": "NON_PARTICIPATING",
    "154": "NON_PARTICIPATING",
    "155": "NON_PARTICIPATING",
    "156": "NOT_ALIVE",
    "157": "NON_PARTICIPATING",
    "158": "NOT_ALIVE",
    "159": "NOT_ALIVE",
    "16": "NON_PARTICIPATING",
    "160": "NON_PARTICIPATING",
    "161": "NON_PARTICIPATING",
    "162": "NOT_ALIVE",
    "163": "NOT_ALIVE",
    "164": "NOT_ALIVE",
    "165": "FAIL",
    "166": "NOT_ALIVE",
    "167": "NOT_ALIVE",
    "168": "NON_PARTICIPATING",
    "169": "NOT_ALIVE",
    "17": "NOT_ALIVE",
    "170": "NON_PARTICIPATING",
    "171": "NOT_ALIVE",
    "172": "NOT_ALIVE",
    "173": "FAIL",
    "174": "NON_PARTICIPATING",
    "175": "NOT_ALIVE",
    "176": "NOT_ALIVE",
    "177": "NOT_ALIVE",
    "178": "NON_PARTICIPATING",
    "179": "NOT_ALIVE",
    "18": "NOT_ALIVE",
    "180": "NOT_ALIVE",
    "181": "NOT_ALIVE",
    "182": "NOT_ALIVE",
    "183": "NOT_ALIVE",
    "184": "NON_PARTICIPATING",
    "185": "NOT_ALIVE",
    "186": "NON_PARTICIPATING",
    "187": "NOT_ALIVE",
    "188": "NON_PARTICIPATING",
    "189": "NOT_ALIVE",
    "19": "NOT_ALIVE",
    "190": "NON_PARTICIPATING",
    "191": "NOT_ALIVE",
    "192": "NOT_ALIVE",
    "193": "NON_PARTICIPATING",
    "194": "NON_PARTICIPATING",
    "195": "NON_PARTICIPATING",
    "196": "NON_PARTICIPATING",
    "197": "NON_PARTICIPATING",
    "198": "NOT_ALIVE",
    "199": "NON_PARTICIPATING",
    "2": "NOT_ALIVE",
    "20": "NOT_ALIVE",
    "200": "NOT_ALIVE",
    "201": "NON_PARTICIPATING",
    "202": "NOT_ALIVE",
    "203": "NOT_ALIVE",
    "204": "SUCCESS",
    "205": "NOT_ALIVE",
    "206": "NON_PARTICIPATING",
    "207": "NOT_ALIVE",
    "208": "NON_PARTICIPATING",
    "209": "NOT_ALIVE",
    "21": "NOT_ALIVE",
    "210": "NON_PARTICIPATING",
    "211": "NON_PARTICIPATING",
    "212": "NON_PARTICIPATING",
    "213": "SUCCESS",
    "214": "NON_PARTICIPATING",
    "215": "NON_PARTICIPATING",
    "216": "NON_PARTICIPATING",
    "217": "NON_PARTICIPATING",
    "218": "NON_PARTICIPATING",
    "219": "NOT_ALIVE",
    "22": "NOT_ALIVE",
    "220": "NON_PARTICIPATING",
    "221": "NOT_ALIVE",
    "222": "NOT_ALIVE",
    "223": "NOT_ALIVE",
    "224": "NOT_ALIVE",
    "225": "NOT_ALIVE",
    "226": "NOT_ALIVE",
    "227": "NON_PARTICIPATING",
    "228": "NON_PARTICIPATING",
    "229": "NON_PARTICIPATING",
    "23": "NON_PARTICIPATING",
    "230": "NOT_ALIVE",
    "231": "NOT_ALIVE",
    "232": "NOT_ALIVE",
    "233": "NOT_ALIVE",
    "234": "NOT_ALIVE",
    "235": "NON_PARTICIPATING",
    "236": "NOT_ALIVE",
    "237": "NON_PARTICIPATING",
    "238": "NON_PARTICIPATING",
    "239": "NOT_ALIVE",
    "24": "NON_PARTICIPATING",
    "240": "NOT_ALIVE",
    "241": "NON_PARTICIPATING",
    "242": "NOT_ALIVE",
    "243": "NOT_ALIVE",
    "244": "NOT_ALIVE",
    "245": "NOT_ALIVE",
    "246": "NON_PARTICIPATING",
    "247": "NON_PARTICIPATING",
    "248": "NOT_ALIVE",
    "249": "NOT_ALIVE",
    "25": "SUCCESS",
    "250": "NON_PARTICIPATING",
    "251": "NOT_ALIVE",
    "252": "NON_PARTICIPATING",
    "253": "NOT_ALIVE",
    "254": "SUCCESS",
    "255": "NON_PARTICIPATING",
    "26": "NON_PARTICIPATING",
    "27": "NON_PARTICIPATING",
    "28": "NOT_ALIVE",
    "29": "NOT_ALIVE",
    "3": "NON_PARTICIPATING",
    "30": "NON_PARTICIPATING",
    "31": "NON_PARTICIPATING",
    "32": "NON_PARTICIPATING",
    "33": "NOT_ALIVE",
    "34": "NOT_ALIVE",
    "35": "SUCCESS",
    "36": "NOT_ALIVE",
    "37": "NOT_ALIVE",
    "38": "NON_PARTICIPATING",
    "39": "FAIL",
    "4": "NOT_ALIVE",
    "40": "NON_PARTICIPATING",
    "41": "NON_PARTICIPATING",
    "42": "NON_PARTICIPATING",
    "43": "NON_PARTICIPATING",
    "44": "NOT_ALIVE",
    "45": "NOT_ALIVE",
    "46": "NOT_ALIVE",
    "47": "NON_PARTICIPATING",
    "48": "NOT_ALIVE",
    "49": "NON_PARTICIPATING",
    "5": "SUCCESS",
    "50": "NOT_ALIVE",
    "51": "NON_PARTICIPATING",
    "52": "SUCCESS",
    "53": "NON_PARTICIPATING",
    "54": "NOT_ALIVE",
    "55": "NOT_ALIVE",
    "56": "NON_PARTICIPATING",
    "57": "NON_PARTICIPATING",
    "58": "NON_PARTICIPATING",
    "59": "SUCCESS",
    "6": "NON_PARTICIPATING",
    "60": "NON_PARTICIPATING",
    "61": "NON_PARTICIPATING",
    "62": "NOT_ALIVE",
    "63": "NON_PARTICIPATING",
    "64": "NOT_ALIVE",
    "65": "NON_PARTICIPATING",
    "66": "NOT_ALIVE",
    "67": "NOT_ALIVE",
    "68": "NON_PARTICIPATING",
    "69": "NOT_ALIVE",
    "7": "NON_PARTICIPATING",
    "70": "NON_PARTICIPATING",
    "71": "NON_PARTICIPATING",
    "72": "NOT_ALIVE",
    "73": "NON_PARTICIPATING",
    "74": "NON_PARTICIPATING",
    "75": "NON_PARTICIPATING",
    "76": "NON_PARTICIPATING",
    "77": "NON_PARTICIPATING",
    "78": "NOT_ALIVE",
    "79": "NOT_ALIVE",
    "8": "NOT_ALIVE",
    "80": "NOT_ALIVE",
    "81": "NOT_ALIVE",
    "82": "NON_PARTICIPATING",
    "83": "NOT_ALIVE",
    "84": "NON_PARTICIPATING",
    "85": "NON_PARTICIPATING",
    "86": "NON_PARTICIPATING",
    "87": "NON_PARTICIPATING",
    "88": "NON_PARTICIPATING",
    "89": "NOT_ALIVE",
    "9": "NON_PARTICIPATING",
    "90": "NON_PARTICIPATING",
    "91": "NON_PARTICIPATING",
    "92": "SUCCESS",
    "93": "NOT_ALIVE",
    "94": "NON_PARTICIPATING",
    "95": "NOT_ALIVE",
    "96": "NON_PARTICIPATING",
    "97": "NOT_ALIVE",
    "98": "NON_PARTICIPATING",
    "99": "NOT_ALIVE"
  },
  "architectures": [
    "LlamaForCausalLM"
  ],
  "attention_bias": false,
  "attention_dropout": 0.0,
  "block_list": [
    5982880,
    5982908
  ],
  "bos_token_id": 1,
  "eos_token_id": 2,
  "hidden_act": "silu",
  "hidden_size": 2048,
  "initializer_range": 0.02,
  "inner_step": 41,
  "intermediate_size": 5632,
  "last_allreduce_block": 5981675,
  "max_position_embeddings": 2048,
  "mlp_bias": false,
  "model_type": "llama",
  "num_attention_heads": 32,
  "num_hidden_layers": 22,
  "num_key_value_heads": 4,
  "pretraining_tp": 1,
  "rms_norm_eps": 1e-05,
  "rope_scaling": null,
  "rope_theta": 10000.0,
  "tie_word_embeddings": false,
  "torch_dtype": "float32",
  "transformers_version": "4.39.3",
  "use_cache": false,
  "vocab_size": 32000
}