Upload ./trainer_log.jsonl with huggingface_hub
Browse files- trainer_log.jsonl +28 -50
trainer_log.jsonl
CHANGED
@@ -1,50 +1,28 @@
|
|
1 |
-
{"current_steps": 1, "total_steps":
|
2 |
-
{"current_steps":
|
3 |
-
{"current_steps":
|
4 |
-
{"current_steps":
|
5 |
-
{"current_steps":
|
6 |
-
{"current_steps":
|
7 |
-
{"current_steps":
|
8 |
-
{"current_steps":
|
9 |
-
{"current_steps":
|
10 |
-
{"current_steps":
|
11 |
-
{"current_steps":
|
12 |
-
{"current_steps":
|
13 |
-
{"current_steps":
|
14 |
-
{"current_steps":
|
15 |
-
{"current_steps":
|
16 |
-
{"current_steps":
|
17 |
-
{"current_steps":
|
18 |
-
{"current_steps":
|
19 |
-
{"current_steps":
|
20 |
-
{"current_steps":
|
21 |
-
{"current_steps":
|
22 |
-
{"current_steps":
|
23 |
-
{"current_steps":
|
24 |
-
{"current_steps":
|
25 |
-
{"current_steps":
|
26 |
-
{"current_steps":
|
27 |
-
{"current_steps":
|
28 |
-
{"current_steps":
|
29 |
-
{"current_steps": 115, "total_steps": 192, "loss": 0.3575, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.329471712759216e-08, "epoch": 2.3958333333333335, "percentage": 59.9, "elapsed_time": "0:09:37", "remaining_time": "0:06:26"}
|
30 |
-
{"current_steps": 120, "total_steps": 192, "loss": 0.3997, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.161995210302015e-08, "epoch": 2.5, "percentage": 62.5, "elapsed_time": "0:09:44", "remaining_time": "0:05:50"}
|
31 |
-
{"current_steps": 120, "total_steps": 192, "loss": null, "eval_loss": 0.8528212904930115, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 2.5, "percentage": 62.5, "elapsed_time": "0:09:44", "remaining_time": "0:05:50"}
|
32 |
-
{"current_steps": 125, "total_steps": 192, "loss": 0.3796, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.075841465580837e-08, "epoch": 2.6041666666666665, "percentage": 65.1, "elapsed_time": "0:11:19", "remaining_time": "0:06:03"}
|
33 |
-
{"current_steps": 130, "total_steps": 192, "loss": 0.3768, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.033564114946932e-08, "epoch": 2.7083333333333335, "percentage": 67.71, "elapsed_time": "0:11:25", "remaining_time": "0:05:26"}
|
34 |
-
{"current_steps": 135, "total_steps": 192, "loss": 0.3762, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.013915282607116e-08, "epoch": 2.8125, "percentage": 70.31, "elapsed_time": "0:11:32", "remaining_time": "0:04:52"}
|
35 |
-
{"current_steps": 140, "total_steps": 192, "loss": 0.3752, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.005343402153039e-08, "epoch": 2.9166666666666665, "percentage": 72.92, "elapsed_time": "0:11:38", "remaining_time": "0:04:19"}
|
36 |
-
{"current_steps": 140, "total_steps": 192, "loss": null, "eval_loss": 0.8572859764099121, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 2.9166666666666665, "percentage": 72.92, "elapsed_time": "0:11:38", "remaining_time": "0:04:19"}
|
37 |
-
{"current_steps": 145, "total_steps": 192, "loss": 0.3768, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.001872829857116e-08, "epoch": 3.0208333333333335, "percentage": 75.52, "elapsed_time": "0:13:12", "remaining_time": "0:04:16"}
|
38 |
-
{"current_steps": 150, "total_steps": 192, "loss": 0.3765, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.000587713853837e-08, "epoch": 3.125, "percentage": 78.12, "elapsed_time": "0:13:21", "remaining_time": "0:03:44"}
|
39 |
-
{"current_steps": 155, "total_steps": 192, "loss": 0.3702, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.0001608748597456e-08, "epoch": 3.2291666666666665, "percentage": 80.73, "elapsed_time": "0:13:28", "remaining_time": "0:03:13"}
|
40 |
-
{"current_steps": 160, "total_steps": 192, "loss": 0.3697, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.0000370319656156e-08, "epoch": 3.3333333333333335, "percentage": 83.33, "elapsed_time": "0:13:35", "remaining_time": "0:02:43"}
|
41 |
-
{"current_steps": 160, "total_steps": 192, "loss": null, "eval_loss": 0.8608274459838867, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 3.3333333333333335, "percentage": 83.33, "elapsed_time": "0:13:35", "remaining_time": "0:02:43"}
|
42 |
-
{"current_steps": 165, "total_steps": 192, "loss": 0.3654, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.0000067945715855e-08, "epoch": 3.4375, "percentage": 85.94, "elapsed_time": "0:15:08", "remaining_time": "0:02:28"}
|
43 |
-
{"current_steps": 170, "total_steps": 192, "loss": 0.3523, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.0000009144677036e-08, "epoch": 3.5416666666666665, "percentage": 88.54, "elapsed_time": "0:15:15", "remaining_time": "0:01:58"}
|
44 |
-
{"current_steps": 175, "total_steps": 192, "loss": 0.3668, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.0000000785521776e-08, "epoch": 3.6458333333333335, "percentage": 91.15, "elapsed_time": "0:15:21", "remaining_time": "0:01:29"}
|
45 |
-
{"current_steps": 180, "total_steps": 192, "loss": 0.3636, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.000000003317662e-08, "epoch": 3.75, "percentage": 93.75, "elapsed_time": "0:15:28", "remaining_time": "0:01:01"}
|
46 |
-
{"current_steps": 180, "total_steps": 192, "loss": null, "eval_loss": 0.8633963465690613, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 3.75, "percentage": 93.75, "elapsed_time": "0:15:28", "remaining_time": "0:01:01"}
|
47 |
-
{"current_steps": 185, "total_steps": 192, "loss": 0.3717, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.000000000038355e-08, "epoch": 3.8541666666666665, "percentage": 96.35, "elapsed_time": "0:17:03", "remaining_time": "0:00:38"}
|
48 |
-
{"current_steps": 190, "total_steps": 192, "loss": 0.3687, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 5.000000000000018e-08, "epoch": 3.9583333333333335, "percentage": 98.96, "elapsed_time": "0:17:10", "remaining_time": "0:00:10"}
|
49 |
-
{"current_steps": 192, "total_steps": 192, "loss": null, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 4.0, "percentage": 100.0, "elapsed_time": "0:17:12", "remaining_time": "0:00:00"}
|
50 |
-
{"current_steps": 3, "total_steps": 3, "loss": null, "eval_loss": 0.8431240916252136, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 4.0, "percentage": 100.0, "elapsed_time": "0:17:45", "remaining_time": "0:00:00"}
|
|
|
1 |
+
{"current_steps": 1, "total_steps": 48, "loss": 1.0289, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 0.0, "epoch": 0.020833333333333332, "percentage": 2.08, "elapsed_time": "0:00:06", "remaining_time": "0:04:54"}
|
2 |
+
{"current_steps": 3, "total_steps": 48, "loss": 0.9475, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 7.5e-07, "epoch": 0.0625, "percentage": 6.25, "elapsed_time": "0:00:07", "remaining_time": "0:01:58"}
|
3 |
+
{"current_steps": 5, "total_steps": 48, "loss": null, "eval_loss": 0.9170231819152832, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.10416666666666667, "percentage": 10.42, "elapsed_time": "0:00:09", "remaining_time": "0:01:20"}
|
4 |
+
{"current_steps": 6, "total_steps": 48, "loss": 1.032, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 3e-06, "epoch": 0.125, "percentage": 12.5, "elapsed_time": "0:00:10", "remaining_time": "0:01:16"}
|
5 |
+
{"current_steps": 9, "total_steps": 48, "loss": 0.8569, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 2.7988636363636366e-06, "epoch": 0.1875, "percentage": 18.75, "elapsed_time": "0:00:13", "remaining_time": "0:00:56"}
|
6 |
+
{"current_steps": 10, "total_steps": 48, "loss": null, "eval_loss": 0.8191882967948914, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.20833333333333334, "percentage": 20.83, "elapsed_time": "0:00:13", "remaining_time": "0:00:51"}
|
7 |
+
{"current_steps": 12, "total_steps": 48, "loss": 0.7857, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 2.6647727272727274e-06, "epoch": 0.25, "percentage": 25.0, "elapsed_time": "0:01:02", "remaining_time": "0:03:08"}
|
8 |
+
{"current_steps": 15, "total_steps": 48, "loss": 0.8358, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 2.4636363636363635e-06, "epoch": 0.3125, "percentage": 31.25, "elapsed_time": "0:01:05", "remaining_time": "0:02:23"}
|
9 |
+
{"current_steps": 15, "total_steps": 48, "loss": null, "eval_loss": 0.8035795092582703, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.3125, "percentage": 31.25, "elapsed_time": "0:01:05", "remaining_time": "0:02:23"}
|
10 |
+
{"current_steps": 18, "total_steps": 48, "loss": 0.7692, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 2.2625e-06, "epoch": 0.375, "percentage": 37.5, "elapsed_time": "0:01:08", "remaining_time": "0:01:53"}
|
11 |
+
{"current_steps": 20, "total_steps": 48, "loss": null, "eval_loss": 0.7883596420288086, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.4166666666666667, "percentage": 41.67, "elapsed_time": "0:01:09", "remaining_time": "0:01:37"}
|
12 |
+
{"current_steps": 21, "total_steps": 48, "loss": 0.8029, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 2.0613636363636364e-06, "epoch": 0.4375, "percentage": 43.75, "elapsed_time": "0:01:59", "remaining_time": "0:02:33"}
|
13 |
+
{"current_steps": 24, "total_steps": 48, "loss": 0.7192, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 1.8602272727272725e-06, "epoch": 0.5, "percentage": 50.0, "elapsed_time": "0:02:01", "remaining_time": "0:02:01"}
|
14 |
+
{"current_steps": 25, "total_steps": 48, "loss": null, "eval_loss": 0.7828572392463684, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.5208333333333334, "percentage": 52.08, "elapsed_time": "0:02:02", "remaining_time": "0:01:52"}
|
15 |
+
{"current_steps": 27, "total_steps": 48, "loss": 0.7906, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 1.659090909090909e-06, "epoch": 0.5625, "percentage": 56.25, "elapsed_time": "0:02:04", "remaining_time": "0:01:36"}
|
16 |
+
{"current_steps": 30, "total_steps": 48, "loss": 0.8181, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 1.4579545454545454e-06, "epoch": 0.625, "percentage": 62.5, "elapsed_time": "0:02:06", "remaining_time": "0:01:15"}
|
17 |
+
{"current_steps": 30, "total_steps": 48, "loss": null, "eval_loss": 0.7780725955963135, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.625, "percentage": 62.5, "elapsed_time": "0:02:06", "remaining_time": "0:01:15"}
|
18 |
+
{"current_steps": 33, "total_steps": 48, "loss": 0.6879, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 1.2568181818181817e-06, "epoch": 0.6875, "percentage": 68.75, "elapsed_time": "0:02:57", "remaining_time": "0:01:20"}
|
19 |
+
{"current_steps": 35, "total_steps": 48, "loss": null, "eval_loss": 0.7747842669487, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.7291666666666666, "percentage": 72.92, "elapsed_time": "0:02:58", "remaining_time": "0:01:06"}
|
20 |
+
{"current_steps": 36, "total_steps": 48, "loss": 0.7508, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 1.0556818181818182e-06, "epoch": 0.75, "percentage": 75.0, "elapsed_time": "0:03:00", "remaining_time": "0:01:00"}
|
21 |
+
{"current_steps": 39, "total_steps": 48, "loss": 0.7032, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 8.545454545454544e-07, "epoch": 0.8125, "percentage": 81.25, "elapsed_time": "0:03:02", "remaining_time": "0:00:42"}
|
22 |
+
{"current_steps": 40, "total_steps": 48, "loss": null, "eval_loss": 0.7724230289459229, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.8333333333333334, "percentage": 83.33, "elapsed_time": "0:03:03", "remaining_time": "0:00:36"}
|
23 |
+
{"current_steps": 42, "total_steps": 48, "loss": 0.7755, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 6.534090909090911e-07, "epoch": 0.875, "percentage": 87.5, "elapsed_time": "0:03:54", "remaining_time": "0:00:33"}
|
24 |
+
{"current_steps": 45, "total_steps": 48, "loss": 0.8136, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 4.522727272727273e-07, "epoch": 0.9375, "percentage": 93.75, "elapsed_time": "0:03:56", "remaining_time": "0:00:15"}
|
25 |
+
{"current_steps": 45, "total_steps": 48, "loss": null, "eval_loss": 0.7702701091766357, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 0.9375, "percentage": 93.75, "elapsed_time": "0:03:56", "remaining_time": "0:00:15"}
|
26 |
+
{"current_steps": 48, "total_steps": 48, "loss": 0.8008, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": 2.511363636363638e-07, "epoch": 1.0, "percentage": 100.0, "elapsed_time": "0:03:59", "remaining_time": "0:00:00"}
|
27 |
+
{"current_steps": 48, "total_steps": 48, "loss": null, "eval_loss": null, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 1.0, "percentage": 100.0, "elapsed_time": "0:03:59", "remaining_time": "0:00:00"}
|
28 |
+
{"current_steps": 3, "total_steps": 3, "loss": null, "eval_loss": 0.7724230289459229, "predict_loss": null, "reward": null, "learning_rate": null, "epoch": 1.0, "percentage": 100.0, "elapsed_time": "0:04:37", "remaining_time": "0:00:00"}
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|
|