Training in progress, step 45000

Browse files

Files changed (10) hide show

last-checkpoint/optimizer.pt +1 -1
last-checkpoint/pytorch_model.bin +1 -1
last-checkpoint/rng_state_0.pth +1 -1
last-checkpoint/rng_state_1.pth +1 -1
last-checkpoint/rng_state_2.pth +1 -1
last-checkpoint/rng_state_3.pth +1 -1
last-checkpoint/scaler.pt +1 -1
last-checkpoint/scheduler.pt +1 -1
last-checkpoint/trainer_state.json +62 -2
pytorch_model.bin +1 -1

last-checkpoint/optimizer.pt CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:350d6a20720d1fd74b8a8444079c82d9172cc3091a8b4988eb3fb998da8e5918
 size 402587859

 version https://git-lfs.github.com/spec/v1
+oid sha256:f7d06d0eceb624a93f15c1ff7d5200ba9ef94286529cf32d4edad055ad078ea0
 size 402587859

last-checkpoint/pytorch_model.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:4f773931870d5ec210c2530416bfc2b566148c206899a79032c3f2237ad80aa1
 size 201355195

 version https://git-lfs.github.com/spec/v1
+oid sha256:03baa03a47a5c1f3c54451145a1853ad32051aa5484ab5c4947f83b3fe9e5f85
 size 201355195

last-checkpoint/rng_state_0.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:8e1d6cdc250e7c41358abfc4778ade95046631664f9a9cd0e1d2397ae0eb957c
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:35ea75c60f66228862cd92e9088324385ce9aed769205f3b1551a63c6fe60f8a
 size 14503

last-checkpoint/rng_state_1.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:ed7ddbb3091468495d36c3725551bda154a5ec0d8b374fc8155ef56a8e8bf136
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:bdf8835b2df6fbbc9c3403b592687a613dd65cafea09e7ef351835386dae4b27
 size 14503

last-checkpoint/rng_state_2.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:d1a882955bcb115cdc4953b3a7056dfd2e484f24ea2b374978a6232b7b1c5d7d
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:474818e6d0125304282978c1c8542265c44eb5493d61b3978a6c3f69535ce59f
 size 14503

last-checkpoint/rng_state_3.pth CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:1307fb0c2ffb9c9e676064788a9a9bfb8168f46f22a121e60cb381f6b8ef60ba
 size 14503

 version https://git-lfs.github.com/spec/v1
+oid sha256:01b5c2f5adb77a01453d635e6c166adf9b4123b6d5fa75b83d11790aeec5a33c
 size 14503

last-checkpoint/scaler.pt CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:418af25aec4bef7120d46157de233e64fab0b4e4ebe06ec15abb44056125efaf
 size 559

 version https://git-lfs.github.com/spec/v1
+oid sha256:bec99c9bd76f872a36dc1adeac79fca3a56a878cf6a77beb9bd3ef47bd234f37
 size 559

last-checkpoint/scheduler.pt CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:a94c54fffa556051788f8ee2aca37f5e903913f8b9cc822ec4876aa0c66a1968
 size 623

 version https://git-lfs.github.com/spec/v1
+oid sha256:b09f8c1cd14309b3eed15dc1e7f72362ffe87586c0b4ccac82d9ab818c4a6f3f
 size 623

last-checkpoint/trainer_state.json CHANGED Viewed

@@ -1,8 +1,8 @@
 {
   "best_metric": null,
   "best_model_checkpoint": null,
-  "epoch": 0.6808510638297872,
-  "global_step": 40000,
   "is_hyper_param_search": false,
   "is_local_process_zero": true,
   "is_world_process_zero": true,
@@ -486,6 +486,66 @@
       "learning_rate": 0.00014780224298665108,
       "loss": 0.3684,
       "step": 40000
     }
   ],
   "max_steps": 500000,

 {
   "best_metric": null,
   "best_model_checkpoint": null,
+  "epoch": 0.7659574468085106,
+  "global_step": 45000,
   "is_hyper_param_search": false,
   "is_local_process_zero": true,
   "is_world_process_zero": true,
       "learning_rate": 0.00014780224298665108,
       "loss": 0.3684,
       "step": 40000
+    },
+    {
+      "epoch": 0.69,
+      "learning_rate": 0.0001477472361221548,
+      "loss": 0.3681,
+      "step": 40500
+    },
+    {
+      "epoch": 0.7,
+      "learning_rate": 0.0001476915606197887,
+      "loss": 0.3677,
+      "step": 41000
+    },
+    {
+      "epoch": 0.71,
+      "learning_rate": 0.00014763521702904748,
+      "loss": 0.3677,
+      "step": 41500
+    },
+    {
+      "epoch": 0.71,
+      "learning_rate": 0.0001475782059060196,
+      "loss": 0.3672,
+      "step": 42000
+    },
+    {
+      "epoch": 0.72,
+      "learning_rate": 0.00014752052781338187,
+      "loss": 0.3673,
+      "step": 42500
+    },
+    {
+      "epoch": 0.73,
+      "learning_rate": 0.00014746218332039373,
+      "loss": 0.3669,
+      "step": 43000
+    },
+    {
+      "epoch": 0.74,
+      "learning_rate": 0.00014740317300289185,
+      "loss": 0.3669,
+      "step": 43500
+    },
+    {
+      "epoch": 0.75,
+      "learning_rate": 0.0001473436174579246,
+      "loss": 0.3666,
+      "step": 44000
+    },
+    {
+      "epoch": 0.76,
+      "learning_rate": 0.00014728327857389895,
+      "loss": 0.3659,
+      "step": 44500
+    },
+    {
+      "epoch": 0.77,
+      "learning_rate": 0.00014722227563107712,
+      "loss": 0.3661,
+      "step": 45000
     }
   ],
   "max_steps": 500000,

pytorch_model.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:4f773931870d5ec210c2530416bfc2b566148c206899a79032c3f2237ad80aa1
 size 201355195

 version https://git-lfs.github.com/spec/v1
+oid sha256:03baa03a47a5c1f3c54451145a1853ad32051aa5484ab5c4947f83b3fe9e5f85
 size 201355195