AmberYifan
commited on
Commit
•
da0d5d8
1
Parent(s):
b9f11e0
Training in progress, step 868, checkpoint
Browse files- last-checkpoint/global_step868/bf16_zero_pp_rank_0_mp_rank_00_optim_states.pt +3 -0
- last-checkpoint/global_step868/bf16_zero_pp_rank_1_mp_rank_00_optim_states.pt +3 -0
- last-checkpoint/global_step868/bf16_zero_pp_rank_2_mp_rank_00_optim_states.pt +3 -0
- last-checkpoint/global_step868/bf16_zero_pp_rank_3_mp_rank_00_optim_states.pt +3 -0
- last-checkpoint/global_step868/zero_pp_rank_0_mp_rank_00_model_states.pt +3 -0
- last-checkpoint/global_step868/zero_pp_rank_1_mp_rank_00_model_states.pt +3 -0
- last-checkpoint/global_step868/zero_pp_rank_2_mp_rank_00_model_states.pt +3 -0
- last-checkpoint/global_step868/zero_pp_rank_3_mp_rank_00_model_states.pt +3 -0
- last-checkpoint/latest +1 -1
- last-checkpoint/model-00001-of-00003.safetensors +1 -1
- last-checkpoint/model-00002-of-00003.safetensors +1 -1
- last-checkpoint/model-00003-of-00003.safetensors +1 -1
- last-checkpoint/rng_state_0.pth +1 -1
- last-checkpoint/rng_state_1.pth +1 -1
- last-checkpoint/rng_state_2.pth +1 -1
- last-checkpoint/rng_state_3.pth +1 -1
- last-checkpoint/scheduler.pt +1 -1
- last-checkpoint/trainer_state.json +108 -2
last-checkpoint/global_step868/bf16_zero_pp_rank_0_mp_rank_00_optim_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:4d373eee01f46747bf9b3ff262ae05ee14ba3f91aed0df0280fdffd6d261d24b
|
3 |
+
size 14483467880
|
last-checkpoint/global_step868/bf16_zero_pp_rank_1_mp_rank_00_optim_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:d08bbefd916d1c9390a1329f91eae56d29b4d4d22ffa16b2dbe169c422e8f610
|
3 |
+
size 14483467880
|
last-checkpoint/global_step868/bf16_zero_pp_rank_2_mp_rank_00_optim_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:3ae39e66019f2289d1d2b61b50b99e0fcdca9934fa72cce6a551f9e5c668e654
|
3 |
+
size 14483467880
|
last-checkpoint/global_step868/bf16_zero_pp_rank_3_mp_rank_00_optim_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:6a51fb19e21a2bf0c4de0b2e8f65e08741e16601bd4ca48d8ba8af00262b9488
|
3 |
+
size 14483467880
|
last-checkpoint/global_step868/zero_pp_rank_0_mp_rank_00_model_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:a3d7e9efb2e141e6e9e6255f7d81521dd9a991493c3a2f561a68005b4b12ea86
|
3 |
+
size 150629
|
last-checkpoint/global_step868/zero_pp_rank_1_mp_rank_00_model_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:9983f08a72661666879082f9f5064686bc2eebea9fb36b48b9cc020c4b74ca7f
|
3 |
+
size 150629
|
last-checkpoint/global_step868/zero_pp_rank_2_mp_rank_00_model_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:f016ebf514bb495af750a075ec720a50234d37eb4d353fd53038f46d22cd2c7e
|
3 |
+
size 150629
|
last-checkpoint/global_step868/zero_pp_rank_3_mp_rank_00_model_states.pt
ADDED
@@ -0,0 +1,3 @@
|
|
|
|
|
|
|
|
|
1 |
+
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:e69702c6c7036fa6579f2b4eb2b267656235a93969516e2add91dc8f26ffa8ba
|
3 |
+
size 150629
|
last-checkpoint/latest
CHANGED
@@ -1 +1 @@
|
|
1 |
-
|
|
|
1 |
+
global_step868
|
last-checkpoint/model-00001-of-00003.safetensors
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 4943162336
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:49ed8f0847a6d218a8d1eccb47868dcceb5eca54475b6157db67d660a1803d1d
|
3 |
size 4943162336
|
last-checkpoint/model-00002-of-00003.safetensors
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 4999819336
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:637ae61efed0ceac33b21a213b84b61d20e3cb817770b79a2eda25413cbb5a3c
|
3 |
size 4999819336
|
last-checkpoint/model-00003-of-00003.safetensors
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 4540516344
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:5d0d7a867e4f5af4d520757d6e400338df292c54d603ec0f25b0254c5ee9e6f5
|
3 |
size 4540516344
|
last-checkpoint/rng_state_0.pth
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 15024
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:8639c02c997d5ec74743bd87a283daff10faa317419bf379edd99c706559f2ce
|
3 |
size 15024
|
last-checkpoint/rng_state_1.pth
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 15024
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:211a79d80fe07a9690b74e693f719eafa8303e6798af58a53dd105eb19c8ccc5
|
3 |
size 15024
|
last-checkpoint/rng_state_2.pth
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 15024
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:319730d0d11be8a12d1633e218e39729160d397a56a34d9ebd2e63d2c81fd68f
|
3 |
size 15024
|
last-checkpoint/rng_state_3.pth
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 15024
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:8e7c1d2c0fa7220ac8b520afb2fc0958467f149187d655695c73de033474c910
|
3 |
size 15024
|
last-checkpoint/scheduler.pt
CHANGED
@@ -1,3 +1,3 @@
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
-
oid sha256:
|
3 |
size 1064
|
|
|
1 |
version https://git-lfs.github.com/spec/v1
|
2 |
+
oid sha256:d0a127df5e8aedd711ac814e5c853ecd07390a0c8ef0dc12756e27ebaa732ecd
|
3 |
size 1064
|
last-checkpoint/trainer_state.json
CHANGED
@@ -1,9 +1,9 @@
|
|
1 |
{
|
2 |
"best_metric": null,
|
3 |
"best_model_checkpoint": null,
|
4 |
-
"epoch": 2.
|
5 |
"eval_steps": 62,
|
6 |
-
"global_step":
|
7 |
"is_hyper_param_search": false,
|
8 |
"is_local_process_zero": true,
|
9 |
"is_world_process_zero": true,
|
@@ -1430,6 +1430,112 @@
|
|
1430 |
"eval_samples_per_second": 5.423,
|
1431 |
"eval_steps_per_second": 0.353,
|
1432 |
"step": 806
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
1433 |
}
|
1434 |
],
|
1435 |
"logging_steps": 10,
|
|
|
1 |
{
|
2 |
"best_metric": null,
|
3 |
"best_model_checkpoint": null,
|
4 |
+
"epoch": 2.7776,
|
5 |
"eval_steps": 62,
|
6 |
+
"global_step": 868,
|
7 |
"is_hyper_param_search": false,
|
8 |
"is_local_process_zero": true,
|
9 |
"is_world_process_zero": true,
|
|
|
1430 |
"eval_samples_per_second": 5.423,
|
1431 |
"eval_steps_per_second": 0.353,
|
1432 |
"step": 806
|
1433 |
+
},
|
1434 |
+
{
|
1435 |
+
"epoch": 2.592,
|
1436 |
+
"grad_norm": 49.93127642652704,
|
1437 |
+
"learning_rate": 7.482185273159145e-08,
|
1438 |
+
"logits/generated": -2.2843379974365234,
|
1439 |
+
"logits/real": -2.4402122497558594,
|
1440 |
+
"logps/generated": -136.36618041992188,
|
1441 |
+
"logps/real": -96.77020263671875,
|
1442 |
+
"loss": 0.2065,
|
1443 |
+
"rewards/accuracies": 0.987500011920929,
|
1444 |
+
"rewards/generated": -0.5222247838973999,
|
1445 |
+
"rewards/margins": 4.2587480545043945,
|
1446 |
+
"rewards/real": 3.736523151397705,
|
1447 |
+
"step": 810
|
1448 |
+
},
|
1449 |
+
{
|
1450 |
+
"epoch": 2.624,
|
1451 |
+
"grad_norm": 37.58250644609933,
|
1452 |
+
"learning_rate": 6.88836104513064e-08,
|
1453 |
+
"logits/generated": -2.3502609729766846,
|
1454 |
+
"logits/real": -2.4047064781188965,
|
1455 |
+
"logps/generated": -141.55758666992188,
|
1456 |
+
"logps/real": -115.49945068359375,
|
1457 |
+
"loss": 0.2793,
|
1458 |
+
"rewards/accuracies": 0.9750000238418579,
|
1459 |
+
"rewards/generated": -0.4805728793144226,
|
1460 |
+
"rewards/margins": 4.331357002258301,
|
1461 |
+
"rewards/real": 3.850783586502075,
|
1462 |
+
"step": 820
|
1463 |
+
},
|
1464 |
+
{
|
1465 |
+
"epoch": 2.656,
|
1466 |
+
"grad_norm": 35.86590095950996,
|
1467 |
+
"learning_rate": 6.294536817102138e-08,
|
1468 |
+
"logits/generated": -2.4404706954956055,
|
1469 |
+
"logits/real": -2.3888516426086426,
|
1470 |
+
"logps/generated": -119.47889709472656,
|
1471 |
+
"logps/real": -102.79316711425781,
|
1472 |
+
"loss": 0.2456,
|
1473 |
+
"rewards/accuracies": 0.949999988079071,
|
1474 |
+
"rewards/generated": -0.43787437677383423,
|
1475 |
+
"rewards/margins": 4.165786266326904,
|
1476 |
+
"rewards/real": 3.7279114723205566,
|
1477 |
+
"step": 830
|
1478 |
+
},
|
1479 |
+
{
|
1480 |
+
"epoch": 2.6879999999999997,
|
1481 |
+
"grad_norm": 72.52101185428741,
|
1482 |
+
"learning_rate": 5.700712589073634e-08,
|
1483 |
+
"logits/generated": -2.3413405418395996,
|
1484 |
+
"logits/real": -2.4445858001708984,
|
1485 |
+
"logps/generated": -143.9036102294922,
|
1486 |
+
"logps/real": -120.9752197265625,
|
1487 |
+
"loss": 0.2809,
|
1488 |
+
"rewards/accuracies": 0.9750000238418579,
|
1489 |
+
"rewards/generated": 0.0636942982673645,
|
1490 |
+
"rewards/margins": 4.376380920410156,
|
1491 |
+
"rewards/real": 4.440074920654297,
|
1492 |
+
"step": 840
|
1493 |
+
},
|
1494 |
+
{
|
1495 |
+
"epoch": 2.7199999999999998,
|
1496 |
+
"grad_norm": 47.87386592895277,
|
1497 |
+
"learning_rate": 5.10688836104513e-08,
|
1498 |
+
"logits/generated": -2.459815740585327,
|
1499 |
+
"logits/real": -2.4833953380584717,
|
1500 |
+
"logps/generated": -145.05397033691406,
|
1501 |
+
"logps/real": -100.8636703491211,
|
1502 |
+
"loss": 0.1964,
|
1503 |
+
"rewards/accuracies": 0.987500011920929,
|
1504 |
+
"rewards/generated": -0.6957341432571411,
|
1505 |
+
"rewards/margins": 4.412557125091553,
|
1506 |
+
"rewards/real": 3.716822385787964,
|
1507 |
+
"step": 850
|
1508 |
+
},
|
1509 |
+
{
|
1510 |
+
"epoch": 2.752,
|
1511 |
+
"grad_norm": 14.759861860387986,
|
1512 |
+
"learning_rate": 4.5130641330166267e-08,
|
1513 |
+
"logits/generated": -2.402906656265259,
|
1514 |
+
"logits/real": -2.3498435020446777,
|
1515 |
+
"logps/generated": -135.8385467529297,
|
1516 |
+
"logps/real": -100.04483795166016,
|
1517 |
+
"loss": 0.2423,
|
1518 |
+
"rewards/accuracies": 0.9750000238418579,
|
1519 |
+
"rewards/generated": -0.7252713441848755,
|
1520 |
+
"rewards/margins": 4.6403326988220215,
|
1521 |
+
"rewards/real": 3.9150619506835938,
|
1522 |
+
"step": 860
|
1523 |
+
},
|
1524 |
+
{
|
1525 |
+
"epoch": 2.7776,
|
1526 |
+
"eval_logits/generated": -2.3574001789093018,
|
1527 |
+
"eval_logits/real": -2.3973894119262695,
|
1528 |
+
"eval_logps/generated": -105.71865844726562,
|
1529 |
+
"eval_logps/real": -116.01789093017578,
|
1530 |
+
"eval_loss": 0.7526271343231201,
|
1531 |
+
"eval_rewards/accuracies": 0.5961538553237915,
|
1532 |
+
"eval_rewards/generated": 1.5596884489059448,
|
1533 |
+
"eval_rewards/margins": 0.7299386262893677,
|
1534 |
+
"eval_rewards/real": 2.2896268367767334,
|
1535 |
+
"eval_runtime": 37.3402,
|
1536 |
+
"eval_samples_per_second": 5.356,
|
1537 |
+
"eval_steps_per_second": 0.348,
|
1538 |
+
"step": 868
|
1539 |
}
|
1540 |
],
|
1541 |
"logging_steps": 10,
|