Saving weights and logs of step 1000
Browse files
events.out.tfevents.1644152162.t1v-n-ccbf3e94-w-0.742851.3.v2
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:f3cc823acc0987cbd5ee0268b216764439077fafaa8a4d0dac6ce85248a66836
|
3 |
+
size 40
|
events.out.tfevents.1644154971.t1v-n-ccbf3e94-w-0.756187.3.v2
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:1ab370b15b8387af7b1f61244d453535d8602bc828029d6095739c6fed2bb9fe
|
3 |
+
size 40
|
events.out.tfevents.1644155872.t1v-n-ccbf3e94-w-0.758395.3.v2
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:cea223d240ece0f2901cb26443d866f4a1e5250ad2b2e0244f21b122b8c573f2
|
3 |
+
size 147136
|
flax_model.msgpack
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 498796983
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:217ab9d795fd2998dddd71aefa7e009066740f1b26675342e3cae64e294b913b
|
3 |
size 498796983
|
run_512.sh
CHANGED
@@ -8,8 +8,8 @@ python run_mlm_flax.py \
|
|
8 |
--cache_dir="/mnt/disks/flaxdisk/cache/" \
|
9 |
--max_seq_length="512" \
|
10 |
--weight_decay="0.01" \
|
11 |
-
--per_device_train_batch_size="
|
12 |
-
--per_device_eval_batch_size="
|
13 |
--pad_to_max_length \
|
14 |
--learning_rate="0.00015" \
|
15 |
--warmup_steps="10000" \
|
|
|
8 |
--cache_dir="/mnt/disks/flaxdisk/cache/" \
|
9 |
--max_seq_length="512" \
|
10 |
--weight_decay="0.01" \
|
11 |
+
--per_device_train_batch_size="46" \
|
12 |
+
--per_device_eval_batch_size="46" \
|
13 |
--pad_to_max_length \
|
14 |
--learning_rate="0.00015" \
|
15 |
--warmup_steps="10000" \
|