Run 2. Outer Step 1. Inner Step 0. Peers 17.
Browse files- config.json +10 -10
- inner_optimizer.pt +1 -1
- model.safetensors +1 -1
- outer_optimizer.pt +1 -1
config.json
CHANGED
@@ -14,7 +14,7 @@
|
|
14 |
"106": "NON_PARTICIPATING",
|
15 |
"107": "NON_PARTICIPATING",
|
16 |
"108": "NON_PARTICIPATING",
|
17 |
-
"109": "
|
18 |
"11": "NON_PARTICIPATING",
|
19 |
"110": "NON_PARTICIPATING",
|
20 |
"111": "NON_PARTICIPATING",
|
@@ -94,10 +94,10 @@
|
|
94 |
"179": "NON_PARTICIPATING",
|
95 |
"18": "NON_PARTICIPATING",
|
96 |
"180": "NON_PARTICIPATING",
|
97 |
-
"181": "
|
98 |
"182": "NON_PARTICIPATING",
|
99 |
"183": "NON_PARTICIPATING",
|
100 |
-
"184": "
|
101 |
"185": "NON_PARTICIPATING",
|
102 |
"186": "NON_PARTICIPATING",
|
103 |
"187": "NON_PARTICIPATING",
|
@@ -106,7 +106,7 @@
|
|
106 |
"19": "NON_PARTICIPATING",
|
107 |
"190": "NON_PARTICIPATING",
|
108 |
"191": "NON_PARTICIPATING",
|
109 |
-
"192": "
|
110 |
"193": "NON_PARTICIPATING",
|
111 |
"194": "NON_PARTICIPATING",
|
112 |
"195": "NON_PARTICIPATING",
|
@@ -114,7 +114,7 @@
|
|
114 |
"197": "NON_PARTICIPATING",
|
115 |
"198": "NON_PARTICIPATING",
|
116 |
"199": "NON_PARTICIPATING",
|
117 |
-
"2": "
|
118 |
"20": "NON_PARTICIPATING",
|
119 |
"200": "SUCCESS",
|
120 |
"201": "NON_PARTICIPATING",
|
@@ -148,13 +148,13 @@
|
|
148 |
"227": "NON_PARTICIPATING",
|
149 |
"228": "NON_PARTICIPATING",
|
150 |
"229": "NON_PARTICIPATING",
|
151 |
-
"23": "
|
152 |
"230": "NON_PARTICIPATING",
|
153 |
"231": "NON_PARTICIPATING",
|
154 |
"232": "NON_PARTICIPATING",
|
155 |
"233": "NON_PARTICIPATING",
|
156 |
"234": "NON_PARTICIPATING",
|
157 |
-
"235": "
|
158 |
"236": "NON_PARTICIPATING",
|
159 |
"237": "NON_PARTICIPATING",
|
160 |
"238": "NON_PARTICIPATING",
|
@@ -212,7 +212,7 @@
|
|
212 |
"55": "NON_PARTICIPATING",
|
213 |
"56": "NON_PARTICIPATING",
|
214 |
"57": "NON_PARTICIPATING",
|
215 |
-
"58": "
|
216 |
"59": "NON_PARTICIPATING",
|
217 |
"6": "NON_PARTICIPATING",
|
218 |
"60": "NON_PARTICIPATING",
|
@@ -234,7 +234,7 @@
|
|
234 |
"75": "NON_PARTICIPATING",
|
235 |
"76": "NON_PARTICIPATING",
|
236 |
"77": "NON_PARTICIPATING",
|
237 |
-
"78": "
|
238 |
"79": "NON_PARTICIPATING",
|
239 |
"8": "NON_PARTICIPATING",
|
240 |
"80": "NON_PARTICIPATING",
|
@@ -275,7 +275,7 @@
|
|
275 |
"initializer_range": 0.02,
|
276 |
"inner_step": 0,
|
277 |
"inner_steps": 0,
|
278 |
-
"last_allreduce_block":
|
279 |
"layer_norm_epsilon": 1e-05,
|
280 |
"model_type": "gpt_optimized",
|
281 |
"n_embd": 1280,
|
|
|
14 |
"106": "NON_PARTICIPATING",
|
15 |
"107": "NON_PARTICIPATING",
|
16 |
"108": "NON_PARTICIPATING",
|
17 |
+
"109": "SUCCESS",
|
18 |
"11": "NON_PARTICIPATING",
|
19 |
"110": "NON_PARTICIPATING",
|
20 |
"111": "NON_PARTICIPATING",
|
|
|
94 |
"179": "NON_PARTICIPATING",
|
95 |
"18": "NON_PARTICIPATING",
|
96 |
"180": "NON_PARTICIPATING",
|
97 |
+
"181": "SUCCESS",
|
98 |
"182": "NON_PARTICIPATING",
|
99 |
"183": "NON_PARTICIPATING",
|
100 |
+
"184": "SUCCESS",
|
101 |
"185": "NON_PARTICIPATING",
|
102 |
"186": "NON_PARTICIPATING",
|
103 |
"187": "NON_PARTICIPATING",
|
|
|
106 |
"19": "NON_PARTICIPATING",
|
107 |
"190": "NON_PARTICIPATING",
|
108 |
"191": "NON_PARTICIPATING",
|
109 |
+
"192": "NON_PARTICIPATING",
|
110 |
"193": "NON_PARTICIPATING",
|
111 |
"194": "NON_PARTICIPATING",
|
112 |
"195": "NON_PARTICIPATING",
|
|
|
114 |
"197": "NON_PARTICIPATING",
|
115 |
"198": "NON_PARTICIPATING",
|
116 |
"199": "NON_PARTICIPATING",
|
117 |
+
"2": "SUCCESS",
|
118 |
"20": "NON_PARTICIPATING",
|
119 |
"200": "SUCCESS",
|
120 |
"201": "NON_PARTICIPATING",
|
|
|
148 |
"227": "NON_PARTICIPATING",
|
149 |
"228": "NON_PARTICIPATING",
|
150 |
"229": "NON_PARTICIPATING",
|
151 |
+
"23": "SUCCESS",
|
152 |
"230": "NON_PARTICIPATING",
|
153 |
"231": "NON_PARTICIPATING",
|
154 |
"232": "NON_PARTICIPATING",
|
155 |
"233": "NON_PARTICIPATING",
|
156 |
"234": "NON_PARTICIPATING",
|
157 |
+
"235": "SUCCESS",
|
158 |
"236": "NON_PARTICIPATING",
|
159 |
"237": "NON_PARTICIPATING",
|
160 |
"238": "NON_PARTICIPATING",
|
|
|
212 |
"55": "NON_PARTICIPATING",
|
213 |
"56": "NON_PARTICIPATING",
|
214 |
"57": "NON_PARTICIPATING",
|
215 |
+
"58": "NON_PARTICIPATING",
|
216 |
"59": "NON_PARTICIPATING",
|
217 |
"6": "NON_PARTICIPATING",
|
218 |
"60": "NON_PARTICIPATING",
|
|
|
234 |
"75": "NON_PARTICIPATING",
|
235 |
"76": "NON_PARTICIPATING",
|
236 |
"77": "NON_PARTICIPATING",
|
237 |
+
"78": "SUCCESS",
|
238 |
"79": "NON_PARTICIPATING",
|
239 |
"8": "NON_PARTICIPATING",
|
240 |
"80": "NON_PARTICIPATING",
|
|
|
275 |
"initializer_range": 0.02,
|
276 |
"inner_step": 0,
|
277 |
"inner_steps": 0,
|
278 |
+
"last_allreduce_block": 5323221,
|
279 |
"layer_norm_epsilon": 1e-05,
|
280 |
"model_type": "gpt_optimized",
|
281 |
"n_embd": 1280,
|
inner_optimizer.pt
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 2752
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:df2588b940f0bab6c967f7abc2a6ff9cd6bad6ed15bd3c1f94527dfc94cc3805
|
3 |
size 2752
|
model.safetensors
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 4040701744
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:72d1f4054db89ccac14b1f1280d753c2392383a6cf8af7ff098a570ffafd3cb9
|
3 |
size 4040701744
|
outer_optimizer.pt
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 4040805354
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1d5f2a8d84f90b6b8ea8f1d13a2aa194614eac91db3c7f307719b7f5a2b485cc
|
3 |
size 4040805354
|