Training in progress, step 10000

Browse files

Files changed (10) hide show

last-checkpoint/optimizer.pt +1 -1
last-checkpoint/pytorch_model.bin +1 -1
last-checkpoint/rng_state_0.pth +1 -1
last-checkpoint/rng_state_1.pth +1 -1
last-checkpoint/rng_state_2.pth +1 -1
last-checkpoint/rng_state_3.pth +1 -1
last-checkpoint/scaler.pt +1 -1
last-checkpoint/scheduler.pt +1 -1
last-checkpoint/trainer_state.json +62 -2
pytorch_model.bin +1 -1

last-checkpoint/optimizer.pt CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:4d9bbdc407c0938638f8e324f185d6daeee392a4bdd014f1725e24909555bcb6
 size 402587859

 version https://git-lfs.github.com/spec/v1
+oid sha256:42b7119dae26f95e759af87d2694035b4627e484829fa085b9887307e584027d
 size 402587859

last-checkpoint/pytorch_model.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:9045a0f15bd76682b21f02d6b8b546976c9deb98b494e4512eb54f69ead13915
 size 201355195

 version https://git-lfs.github.com/spec/v1
+oid sha256:0d721a5d92de7c3e7b282220f9260ab6283761aa6c0a67633f01b73f4484e88a
 size 201355195

last-checkpoint/rng_state_0.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:2f7f70c5c76d6ebed5e2ebb52ee49f2e9751c887c300446944a9b31f8021a9ee
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:ced61c2ca16a8b9af34ec22215625eca8ece206506bd667270651e29c50d152c
 size 14503

last-checkpoint/rng_state_1.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:3b681bc631cfe338fd6f57e18faef57f88c46490dd2896345dceffec93f9d383
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:8fab9432af2478fa523c3f480312c9b289cccc342f6c08d159285299a420f2d2
 size 14503

last-checkpoint/rng_state_2.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:4aed1ee46e1319cf55dec13d95fdf2a40998c4089d9cfe71fb6e5127cf2a6158
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:8de7b9d2730e6e1ffa628554e02f762aed35f57534bbe4eca6576f8bf323a2b8
 size 14503

last-checkpoint/rng_state_3.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:8ca0d5c8d530eccaf75abaec019ac769cb22903ecbb4731f2e6ad5a04a9dad19
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:351b749df84f2ebffd7b8765a4b4f3c28b010ec954ada54a535dd8e6072b32f3
 size 14503

last-checkpoint/scaler.pt CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:018632b7efb97680221078f777b72fc854385a42730f0e76446b2116c7b5ff33
 size 559

 version https://git-lfs.github.com/spec/v1
+oid sha256:2a28e8be2a3887ca8072f99cc9bcae55040769870a8c9d9965493ca683763af7
 size 559

last-checkpoint/scheduler.pt CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:d0e8ba5e8a79631162151d1ac9f60167dfa7fb6dea5ba86993956bb6b1d9e50b
 size 623

 version https://git-lfs.github.com/spec/v1
+oid sha256:a3134a72416abec950fef9303930c8b660ea9e82e18cbbb0f68ee969677327d0
 size 623

last-checkpoint/trainer_state.json CHANGED Viewed

@@ -1,8 +1,8 @@
 {
   "best_metric": null,
   "best_model_checkpoint": null,
-  "epoch": 0.08510565867524532,
-  "global_step": 5000,
   "is_hyper_param_search": false,
   "is_local_process_zero": true,
   "is_world_process_zero": true,
@@ -66,6 +66,66 @@
       "learning_rate": 0.0001499654592256012,
       "loss": 0.4466,
       "step": 5000
     }
   ],
   "max_steps": 500000,

 {
   "best_metric": null,
   "best_model_checkpoint": null,
+  "epoch": 0.17021131735049064,
+  "global_step": 10000,
   "is_hyper_param_search": false,
   "is_local_process_zero": true,
   "is_world_process_zero": true,
       "learning_rate": 0.0001499654592256012,
       "loss": 0.4466,
       "step": 5000
+    },
+    {
+      "epoch": 0.09,
+      "learning_rate": 0.0001499582063848481,
+      "loss": 0.4398,
+      "step": 5500
+    },
+    {
+      "epoch": 0.1,
+      "learning_rate": 0.00014995026308484122,
+      "loss": 0.4337,
+      "step": 6000
+    },
+    {
+      "epoch": 0.11,
+      "learning_rate": 0.00014994162940397778,
+      "loss": 0.4286,
+      "step": 6500
+    },
+    {
+      "epoch": 0.12,
+      "learning_rate": 0.00014993230542746872,
+      "loss": 0.4242,
+      "step": 7000
+    },
+    {
+      "epoch": 0.13,
+      "learning_rate": 0.00014992229124733788,
+      "loss": 0.4203,
+      "step": 7500
+    },
+    {
+      "epoch": 0.14,
+      "learning_rate": 0.0001499115869624212,
+      "loss": 0.4173,
+      "step": 8000
+    },
+    {
+      "epoch": 0.14,
+      "learning_rate": 0.00014990019267836565,
+      "loss": 0.4141,
+      "step": 8500
+    },
+    {
+      "epoch": 0.15,
+      "learning_rate": 0.00014988810850762824,
+      "loss": 0.4116,
+      "step": 9000
+    },
+    {
+      "epoch": 0.16,
+      "learning_rate": 0.00014987533456947482,
+      "loss": 0.4095,
+      "step": 9500
+    },
+    {
+      "epoch": 0.17,
+      "learning_rate": 0.000149861870989979,
+      "loss": 0.4075,
+      "step": 10000
     }
   ],
   "max_steps": 500000,

pytorch_model.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:9045a0f15bd76682b21f02d6b8b546976c9deb98b494e4512eb54f69ead13915
 size 201355195

 version https://git-lfs.github.com/spec/v1
+oid sha256:0d721a5d92de7c3e7b282220f9260ab6283761aa6c0a67633f01b73f4484e88a
 size 201355195