llama2_inst_truthbench1_model / trainer_log.jsonl
Ogamon's picture
Initial commit
62c989d verified
raw
history blame
48 kB
{"current_steps": 1, "total_steps": 190, "loss": 8.3599, "learning_rate": 5.000000000000001e-07, "epoch": 0.02572347266881029, "percentage": 0.53, "elapsed_time": "0:00:12", "remaining_time": "0:40:53", "throughput": "548.54", "total_tokens": 7120}
{"current_steps": 2, "total_steps": 190, "loss": 8.1891, "learning_rate": 1.0000000000000002e-06, "epoch": 0.05144694533762058, "percentage": 1.05, "elapsed_time": "0:00:24", "remaining_time": "0:37:46", "throughput": "575.99", "total_tokens": 13888}
{"current_steps": 3, "total_steps": 190, "loss": 8.0792, "learning_rate": 1.5e-06, "epoch": 0.07717041800643087, "percentage": 1.58, "elapsed_time": "0:00:35", "remaining_time": "0:36:35", "throughput": "586.40", "total_tokens": 20656}
{"current_steps": 4, "total_steps": 190, "loss": 7.9682, "learning_rate": 2.0000000000000003e-06, "epoch": 0.10289389067524116, "percentage": 2.11, "elapsed_time": "0:00:46", "remaining_time": "0:35:53", "throughput": "586.87", "total_tokens": 27184}
{"current_steps": 5, "total_steps": 190, "loss": 6.9482, "learning_rate": 2.5e-06, "epoch": 0.12861736334405144, "percentage": 2.63, "elapsed_time": "0:00:57", "remaining_time": "0:35:24", "throughput": "599.42", "total_tokens": 34416}
{"current_steps": 6, "total_steps": 190, "loss": 5.1505, "learning_rate": 3e-06, "epoch": 0.15434083601286175, "percentage": 3.16, "elapsed_time": "0:01:08", "remaining_time": "0:35:00", "throughput": "599.28", "total_tokens": 41056}
{"current_steps": 7, "total_steps": 190, "loss": 4.7491, "learning_rate": 3.5e-06, "epoch": 0.18006430868167203, "percentage": 3.68, "elapsed_time": "0:01:19", "remaining_time": "0:34:41", "throughput": "596.99", "total_tokens": 47536}
{"current_steps": 8, "total_steps": 190, "loss": 3.2164, "learning_rate": 4.000000000000001e-06, "epoch": 0.2057877813504823, "percentage": 4.21, "elapsed_time": "0:01:30", "remaining_time": "0:34:24", "throughput": "600.21", "total_tokens": 54464}
{"current_steps": 9, "total_steps": 190, "loss": 2.7761, "learning_rate": 4.5e-06, "epoch": 0.2315112540192926, "percentage": 4.74, "elapsed_time": "0:01:41", "remaining_time": "0:34:08", "throughput": "603.94", "total_tokens": 61520}
{"current_steps": 10, "total_steps": 190, "loss": 0.6703, "learning_rate": 5e-06, "epoch": 0.2572347266881029, "percentage": 5.26, "elapsed_time": "0:01:52", "remaining_time": "0:33:53", "throughput": "605.96", "total_tokens": 68464}
{"current_steps": 11, "total_steps": 190, "loss": 0.3255, "learning_rate": 4.9996192378909785e-06, "epoch": 0.2829581993569132, "percentage": 5.79, "elapsed_time": "0:02:04", "remaining_time": "0:33:39", "throughput": "605.48", "total_tokens": 75152}
{"current_steps": 12, "total_steps": 190, "loss": 0.3301, "learning_rate": 4.99847706754774e-06, "epoch": 0.3086816720257235, "percentage": 6.32, "elapsed_time": "0:02:15", "remaining_time": "0:33:25", "throughput": "605.64", "total_tokens": 81888}
{"current_steps": 13, "total_steps": 190, "loss": 0.2121, "learning_rate": 4.9965738368864345e-06, "epoch": 0.33440514469453375, "percentage": 6.84, "elapsed_time": "0:02:26", "remaining_time": "0:33:12", "throughput": "606.15", "total_tokens": 88688}
{"current_steps": 14, "total_steps": 190, "loss": 1.1565, "learning_rate": 4.993910125649561e-06, "epoch": 0.36012861736334406, "percentage": 7.37, "elapsed_time": "0:02:37", "remaining_time": "0:32:58", "throughput": "607.41", "total_tokens": 95616}
{"current_steps": 15, "total_steps": 190, "loss": 0.8054, "learning_rate": 4.990486745229364e-06, "epoch": 0.3858520900321543, "percentage": 7.89, "elapsed_time": "0:02:48", "remaining_time": "0:32:45", "throughput": "606.19", "total_tokens": 102144}
{"current_steps": 16, "total_steps": 190, "loss": 0.2386, "learning_rate": 4.986304738420684e-06, "epoch": 0.4115755627009646, "percentage": 8.42, "elapsed_time": "0:02:59", "remaining_time": "0:32:33", "throughput": "607.52", "total_tokens": 109120}
{"current_steps": 17, "total_steps": 190, "loss": 0.3161, "learning_rate": 4.981365379103306e-06, "epoch": 0.43729903536977494, "percentage": 8.95, "elapsed_time": "0:03:10", "remaining_time": "0:32:21", "throughput": "606.78", "total_tokens": 115744}
{"current_steps": 18, "total_steps": 190, "loss": 0.2773, "learning_rate": 4.975670171853926e-06, "epoch": 0.4630225080385852, "percentage": 9.47, "elapsed_time": "0:03:21", "remaining_time": "0:32:09", "throughput": "607.80", "total_tokens": 122704}
{"current_steps": 19, "total_steps": 190, "loss": 0.2062, "learning_rate": 4.9692208514878445e-06, "epoch": 0.4887459807073955, "percentage": 10.0, "elapsed_time": "0:03:33", "remaining_time": "0:31:57", "throughput": "608.19", "total_tokens": 129552}
{"current_steps": 20, "total_steps": 190, "loss": 0.1837, "learning_rate": 4.962019382530521e-06, "epoch": 0.5144694533762058, "percentage": 10.53, "elapsed_time": "0:03:44", "remaining_time": "0:31:45", "throughput": "609.22", "total_tokens": 136544}
{"current_steps": 21, "total_steps": 190, "loss": 0.1735, "learning_rate": 4.9540679586191605e-06, "epoch": 0.5401929260450161, "percentage": 11.05, "elapsed_time": "0:03:55", "remaining_time": "0:31:32", "throughput": "610.01", "total_tokens": 143488}
{"current_steps": 22, "total_steps": 190, "loss": 0.1588, "learning_rate": 4.9453690018345144e-06, "epoch": 0.5659163987138264, "percentage": 11.58, "elapsed_time": "0:04:06", "remaining_time": "0:31:20", "throughput": "609.91", "total_tokens": 150224}
{"current_steps": 23, "total_steps": 190, "loss": 0.1443, "learning_rate": 4.935925161963089e-06, "epoch": 0.5916398713826366, "percentage": 12.11, "elapsed_time": "0:04:17", "remaining_time": "0:31:09", "throughput": "610.82", "total_tokens": 157232}
{"current_steps": 24, "total_steps": 190, "loss": 0.157, "learning_rate": 4.925739315689991e-06, "epoch": 0.617363344051447, "percentage": 12.63, "elapsed_time": "0:04:28", "remaining_time": "0:30:57", "throughput": "609.97", "total_tokens": 163776}
{"current_steps": 25, "total_steps": 190, "loss": 0.1199, "learning_rate": 4.914814565722671e-06, "epoch": 0.6430868167202572, "percentage": 13.16, "elapsed_time": "0:04:39", "remaining_time": "0:30:45", "throughput": "609.21", "total_tokens": 170352}
{"current_steps": 26, "total_steps": 190, "loss": 0.1539, "learning_rate": 4.903154239845798e-06, "epoch": 0.6688102893890675, "percentage": 13.68, "elapsed_time": "0:04:50", "remaining_time": "0:30:33", "throughput": "609.20", "total_tokens": 177120}
{"current_steps": 27, "total_steps": 190, "loss": 0.1208, "learning_rate": 4.890761889907589e-06, "epoch": 0.6945337620578779, "percentage": 14.21, "elapsed_time": "0:05:01", "remaining_time": "0:30:22", "throughput": "609.87", "total_tokens": 184096}
{"current_steps": 28, "total_steps": 190, "loss": 0.0954, "learning_rate": 4.8776412907378845e-06, "epoch": 0.7202572347266881, "percentage": 14.74, "elapsed_time": "0:05:12", "remaining_time": "0:30:10", "throughput": "610.39", "total_tokens": 191040}
{"current_steps": 29, "total_steps": 190, "loss": 0.1387, "learning_rate": 4.863796438998293e-06, "epoch": 0.7459807073954984, "percentage": 15.26, "elapsed_time": "0:05:24", "remaining_time": "0:29:59", "throughput": "611.13", "total_tokens": 198064}
{"current_steps": 30, "total_steps": 190, "loss": 0.1484, "learning_rate": 4.849231551964771e-06, "epoch": 0.7717041800643086, "percentage": 15.79, "elapsed_time": "0:05:35", "remaining_time": "0:29:47", "throughput": "612.02", "total_tokens": 205136}
{"current_steps": 31, "total_steps": 190, "loss": 0.0998, "learning_rate": 4.833951066243004e-06, "epoch": 0.797427652733119, "percentage": 16.32, "elapsed_time": "0:05:46", "remaining_time": "0:29:36", "throughput": "612.22", "total_tokens": 212000}
{"current_steps": 32, "total_steps": 190, "loss": 0.1068, "learning_rate": 4.817959636416969e-06, "epoch": 0.8231511254019293, "percentage": 16.84, "elapsed_time": "0:05:57", "remaining_time": "0:29:24", "throughput": "612.05", "total_tokens": 218720}
{"current_steps": 33, "total_steps": 190, "loss": 0.0801, "learning_rate": 4.801262133631101e-06, "epoch": 0.8488745980707395, "percentage": 17.37, "elapsed_time": "0:06:08", "remaining_time": "0:29:12", "throughput": "612.99", "total_tokens": 225856}
{"current_steps": 34, "total_steps": 190, "loss": 0.1066, "learning_rate": 4.783863644106502e-06, "epoch": 0.8745980707395499, "percentage": 17.89, "elapsed_time": "0:06:19", "remaining_time": "0:29:01", "throughput": "612.89", "total_tokens": 232640}
{"current_steps": 35, "total_steps": 190, "loss": 0.1038, "learning_rate": 4.765769467591626e-06, "epoch": 0.9003215434083601, "percentage": 18.42, "elapsed_time": "0:06:30", "remaining_time": "0:28:50", "throughput": "613.01", "total_tokens": 239504}
{"current_steps": 36, "total_steps": 190, "loss": 0.106, "learning_rate": 4.746985115747918e-06, "epoch": 0.9260450160771704, "percentage": 18.95, "elapsed_time": "0:06:41", "remaining_time": "0:28:38", "throughput": "612.94", "total_tokens": 246288}
{"current_steps": 37, "total_steps": 190, "loss": 0.1107, "learning_rate": 4.72751631047092e-06, "epoch": 0.9517684887459807, "percentage": 19.47, "elapsed_time": "0:06:52", "remaining_time": "0:28:27", "throughput": "613.01", "total_tokens": 253136}
{"current_steps": 38, "total_steps": 190, "loss": 0.1372, "learning_rate": 4.707368982147318e-06, "epoch": 0.977491961414791, "percentage": 20.0, "elapsed_time": "0:07:04", "remaining_time": "0:28:16", "throughput": "613.54", "total_tokens": 260160}
{"current_steps": 39, "total_steps": 190, "loss": 0.0816, "learning_rate": 4.68654926784849e-06, "epoch": 1.0032154340836013, "percentage": 20.53, "elapsed_time": "0:07:15", "remaining_time": "0:28:04", "throughput": "613.88", "total_tokens": 267120}
{"current_steps": 40, "total_steps": 190, "loss": 0.0743, "learning_rate": 4.665063509461098e-06, "epoch": 1.0289389067524115, "percentage": 21.05, "elapsed_time": "0:07:26", "remaining_time": "0:27:53", "throughput": "614.30", "total_tokens": 274112}
{"current_steps": 41, "total_steps": 190, "loss": 0.072, "learning_rate": 4.642918251755281e-06, "epoch": 1.0546623794212218, "percentage": 21.58, "elapsed_time": "0:07:37", "remaining_time": "0:27:41", "throughput": "614.77", "total_tokens": 281136}
{"current_steps": 42, "total_steps": 190, "loss": 0.0596, "learning_rate": 4.620120240391065e-06, "epoch": 1.0803858520900322, "percentage": 22.11, "elapsed_time": "0:07:48", "remaining_time": "0:27:30", "throughput": "614.97", "total_tokens": 288048}
{"current_steps": 43, "total_steps": 190, "loss": 0.0544, "learning_rate": 4.596676419863561e-06, "epoch": 1.1061093247588425, "percentage": 22.63, "elapsed_time": "0:07:59", "remaining_time": "0:27:19", "throughput": "615.46", "total_tokens": 295120}
{"current_steps": 44, "total_steps": 190, "loss": 0.0342, "learning_rate": 4.572593931387604e-06, "epoch": 1.1318327974276527, "percentage": 23.16, "elapsed_time": "0:08:10", "remaining_time": "0:27:07", "throughput": "615.55", "total_tokens": 302000}
{"current_steps": 45, "total_steps": 190, "loss": 0.0394, "learning_rate": 4.54788011072248e-06, "epoch": 1.157556270096463, "percentage": 23.68, "elapsed_time": "0:08:21", "remaining_time": "0:26:56", "throughput": "615.19", "total_tokens": 308672}
{"current_steps": 46, "total_steps": 190, "loss": 0.0196, "learning_rate": 4.522542485937369e-06, "epoch": 1.1832797427652733, "percentage": 24.21, "elapsed_time": "0:08:32", "remaining_time": "0:26:45", "throughput": "615.36", "total_tokens": 315600}
{"current_steps": 47, "total_steps": 190, "loss": 0.0411, "learning_rate": 4.496588775118232e-06, "epoch": 1.2090032154340835, "percentage": 24.74, "elapsed_time": "0:08:43", "remaining_time": "0:26:34", "throughput": "615.43", "total_tokens": 322464}
{"current_steps": 48, "total_steps": 190, "loss": 0.0257, "learning_rate": 4.470026884016805e-06, "epoch": 1.234726688102894, "percentage": 25.26, "elapsed_time": "0:08:55", "remaining_time": "0:26:22", "throughput": "614.94", "total_tokens": 329024}
{"current_steps": 49, "total_steps": 190, "loss": 0.0289, "learning_rate": 4.442864903642428e-06, "epoch": 1.2604501607717042, "percentage": 25.79, "elapsed_time": "0:09:06", "remaining_time": "0:26:11", "throughput": "615.29", "total_tokens": 336032}
{"current_steps": 50, "total_steps": 190, "loss": 0.1193, "learning_rate": 4.415111107797445e-06, "epoch": 1.2861736334405145, "percentage": 26.32, "elapsed_time": "0:09:17", "remaining_time": "0:26:00", "throughput": "615.01", "total_tokens": 342704}
{"current_steps": 51, "total_steps": 190, "loss": 0.0883, "learning_rate": 4.386773950556931e-06, "epoch": 1.3118971061093248, "percentage": 26.84, "elapsed_time": "0:09:28", "remaining_time": "0:25:48", "throughput": "614.92", "total_tokens": 349472}
{"current_steps": 52, "total_steps": 190, "loss": 0.0377, "learning_rate": 4.357862063693486e-06, "epoch": 1.337620578778135, "percentage": 27.37, "elapsed_time": "0:09:39", "remaining_time": "0:25:37", "throughput": "614.86", "total_tokens": 356272}
{"current_steps": 53, "total_steps": 190, "loss": 0.0602, "learning_rate": 4.328384254047927e-06, "epoch": 1.3633440514469453, "percentage": 27.89, "elapsed_time": "0:09:50", "remaining_time": "0:25:26", "throughput": "614.73", "total_tokens": 363040}
{"current_steps": 54, "total_steps": 190, "loss": 0.083, "learning_rate": 4.2983495008466285e-06, "epoch": 1.3890675241157555, "percentage": 28.42, "elapsed_time": "0:10:01", "remaining_time": "0:25:15", "throughput": "614.38", "total_tokens": 369664}
{"current_steps": 55, "total_steps": 190, "loss": 0.0358, "learning_rate": 4.267766952966369e-06, "epoch": 1.414790996784566, "percentage": 28.95, "elapsed_time": "0:10:12", "remaining_time": "0:25:04", "throughput": "614.72", "total_tokens": 376704}
{"current_steps": 56, "total_steps": 190, "loss": 0.0321, "learning_rate": 4.236645926147493e-06, "epoch": 1.4405144694533762, "percentage": 29.47, "elapsed_time": "0:10:23", "remaining_time": "0:24:52", "throughput": "614.84", "total_tokens": 383600}
{"current_steps": 57, "total_steps": 190, "loss": 0.0452, "learning_rate": 4.204995900156247e-06, "epoch": 1.4662379421221865, "percentage": 30.0, "elapsed_time": "0:10:34", "remaining_time": "0:24:41", "throughput": "615.11", "total_tokens": 390592}
{"current_steps": 58, "total_steps": 190, "loss": 0.0915, "learning_rate": 4.172826515897146e-06, "epoch": 1.4919614147909968, "percentage": 30.53, "elapsed_time": "0:10:46", "remaining_time": "0:24:30", "throughput": "615.02", "total_tokens": 397360}
{"current_steps": 59, "total_steps": 190, "loss": 0.0651, "learning_rate": 4.140147572476269e-06, "epoch": 1.517684887459807, "percentage": 31.05, "elapsed_time": "0:10:57", "remaining_time": "0:24:19", "throughput": "614.81", "total_tokens": 404048}
{"current_steps": 60, "total_steps": 190, "loss": 0.0868, "learning_rate": 4.106969024216348e-06, "epoch": 1.5434083601286175, "percentage": 31.58, "elapsed_time": "0:11:08", "remaining_time": "0:24:08", "throughput": "614.92", "total_tokens": 410960}
{"current_steps": 61, "total_steps": 190, "loss": 0.0554, "learning_rate": 4.073300977624594e-06, "epoch": 1.5691318327974275, "percentage": 32.11, "elapsed_time": "0:11:19", "remaining_time": "0:23:56", "throughput": "615.06", "total_tokens": 417888}
{"current_steps": 62, "total_steps": 190, "loss": 0.0336, "learning_rate": 4.039153688314146e-06, "epoch": 1.594855305466238, "percentage": 32.63, "elapsed_time": "0:11:30", "remaining_time": "0:23:45", "throughput": "615.29", "total_tokens": 424880}
{"current_steps": 63, "total_steps": 190, "loss": 0.0455, "learning_rate": 4.0045375578801216e-06, "epoch": 1.6205787781350482, "percentage": 33.16, "elapsed_time": "0:11:41", "remaining_time": "0:23:34", "throughput": "615.69", "total_tokens": 432000}
{"current_steps": 64, "total_steps": 190, "loss": 0.0406, "learning_rate": 3.969463130731183e-06, "epoch": 1.6463022508038585, "percentage": 33.68, "elapsed_time": "0:11:52", "remaining_time": "0:23:23", "throughput": "615.45", "total_tokens": 438672}
{"current_steps": 65, "total_steps": 190, "loss": 0.0461, "learning_rate": 3.933941090877615e-06, "epoch": 1.6720257234726688, "percentage": 34.21, "elapsed_time": "0:12:03", "remaining_time": "0:23:12", "throughput": "615.37", "total_tokens": 445440}
{"current_steps": 66, "total_steps": 190, "loss": 0.0466, "learning_rate": 3.897982258676867e-06, "epoch": 1.697749196141479, "percentage": 34.74, "elapsed_time": "0:12:14", "remaining_time": "0:23:00", "throughput": "615.10", "total_tokens": 452064}
{"current_steps": 67, "total_steps": 190, "loss": 0.0382, "learning_rate": 3.861597587537568e-06, "epoch": 1.7234726688102895, "percentage": 35.26, "elapsed_time": "0:12:26", "remaining_time": "0:22:49", "throughput": "615.23", "total_tokens": 458992}
{"current_steps": 68, "total_steps": 190, "loss": 0.0426, "learning_rate": 3.824798160583012e-06, "epoch": 1.7491961414790995, "percentage": 35.79, "elapsed_time": "0:12:37", "remaining_time": "0:22:38", "throughput": "614.90", "total_tokens": 465568}
{"current_steps": 69, "total_steps": 190, "loss": 0.0264, "learning_rate": 3.787595187275136e-06, "epoch": 1.77491961414791, "percentage": 36.32, "elapsed_time": "0:12:48", "remaining_time": "0:22:27", "throughput": "615.03", "total_tokens": 472496}
{"current_steps": 70, "total_steps": 190, "loss": 0.0567, "learning_rate": 3.7500000000000005e-06, "epoch": 1.8006430868167203, "percentage": 36.84, "elapsed_time": "0:12:59", "remaining_time": "0:22:16", "throughput": "615.11", "total_tokens": 479392}
{"current_steps": 71, "total_steps": 190, "loss": 0.0688, "learning_rate": 3.7120240506158433e-06, "epoch": 1.8263665594855305, "percentage": 37.37, "elapsed_time": "0:13:10", "remaining_time": "0:22:04", "throughput": "615.35", "total_tokens": 486416}
{"current_steps": 72, "total_steps": 190, "loss": 0.0351, "learning_rate": 3.6736789069647273e-06, "epoch": 1.852090032154341, "percentage": 37.89, "elapsed_time": "0:13:21", "remaining_time": "0:21:53", "throughput": "614.88", "total_tokens": 492896}
{"current_steps": 73, "total_steps": 190, "loss": 0.0246, "learning_rate": 3.634976249348867e-06, "epoch": 1.877813504823151, "percentage": 38.42, "elapsed_time": "0:13:32", "remaining_time": "0:21:42", "throughput": "614.93", "total_tokens": 499760}
{"current_steps": 74, "total_steps": 190, "loss": 0.0364, "learning_rate": 3.595927866972694e-06, "epoch": 1.9035369774919615, "percentage": 38.95, "elapsed_time": "0:13:43", "remaining_time": "0:21:31", "throughput": "615.23", "total_tokens": 506816}
{"current_steps": 75, "total_steps": 190, "loss": 0.0352, "learning_rate": 3.556545654351749e-06, "epoch": 1.9292604501607717, "percentage": 39.47, "elapsed_time": "0:13:54", "remaining_time": "0:21:20", "throughput": "615.23", "total_tokens": 513648}
{"current_steps": 76, "total_steps": 190, "loss": 0.0915, "learning_rate": 3.516841607689501e-06, "epoch": 1.954983922829582, "percentage": 40.0, "elapsed_time": "0:14:05", "remaining_time": "0:21:08", "throughput": "615.24", "total_tokens": 520480}
{"current_steps": 77, "total_steps": 190, "loss": 0.0327, "learning_rate": 3.476827821223184e-06, "epoch": 1.9807073954983923, "percentage": 40.53, "elapsed_time": "0:14:17", "remaining_time": "0:20:57", "throughput": "614.95", "total_tokens": 527056}
{"current_steps": 78, "total_steps": 190, "loss": 0.0448, "learning_rate": 3.436516483539781e-06, "epoch": 2.0064308681672025, "percentage": 41.05, "elapsed_time": "0:14:28", "remaining_time": "0:20:46", "throughput": "615.21", "total_tokens": 534112}
{"current_steps": 79, "total_steps": 190, "loss": 0.0186, "learning_rate": 3.39591987386325e-06, "epoch": 2.032154340836013, "percentage": 41.58, "elapsed_time": "0:14:39", "remaining_time": "0:20:35", "throughput": "615.29", "total_tokens": 541024}
{"current_steps": 80, "total_steps": 190, "loss": 0.0342, "learning_rate": 3.3550503583141726e-06, "epoch": 2.057877813504823, "percentage": 42.11, "elapsed_time": "0:14:50", "remaining_time": "0:20:24", "throughput": "615.30", "total_tokens": 547888}
{"current_steps": 81, "total_steps": 190, "loss": 0.0079, "learning_rate": 3.313920386142892e-06, "epoch": 2.0836012861736335, "percentage": 42.63, "elapsed_time": "0:15:01", "remaining_time": "0:20:13", "throughput": "615.14", "total_tokens": 554592}
{"current_steps": 82, "total_steps": 190, "loss": 0.0177, "learning_rate": 3.272542485937369e-06, "epoch": 2.1093247588424435, "percentage": 43.16, "elapsed_time": "0:15:12", "remaining_time": "0:20:02", "throughput": "615.01", "total_tokens": 561296}
{"current_steps": 83, "total_steps": 190, "loss": 0.0139, "learning_rate": 3.230929261806842e-06, "epoch": 2.135048231511254, "percentage": 43.68, "elapsed_time": "0:15:23", "remaining_time": "0:19:50", "throughput": "614.74", "total_tokens": 567872}
{"current_steps": 84, "total_steps": 190, "loss": 0.0103, "learning_rate": 3.189093389542498e-06, "epoch": 2.1607717041800645, "percentage": 44.21, "elapsed_time": "0:15:34", "remaining_time": "0:19:39", "throughput": "615.15", "total_tokens": 575072}
{"current_steps": 85, "total_steps": 190, "loss": 0.0221, "learning_rate": 3.147047612756302e-06, "epoch": 2.1864951768488745, "percentage": 44.74, "elapsed_time": "0:15:45", "remaining_time": "0:19:28", "throughput": "615.35", "total_tokens": 582080}
{"current_steps": 86, "total_steps": 190, "loss": 0.0021, "learning_rate": 3.1048047389991693e-06, "epoch": 2.212218649517685, "percentage": 45.26, "elapsed_time": "0:15:57", "remaining_time": "0:19:17", "throughput": "615.26", "total_tokens": 588816}
{"current_steps": 87, "total_steps": 190, "loss": 0.011, "learning_rate": 3.062377635859663e-06, "epoch": 2.237942122186495, "percentage": 45.79, "elapsed_time": "0:16:08", "remaining_time": "0:19:06", "throughput": "615.65", "total_tokens": 596032}
{"current_steps": 88, "total_steps": 190, "loss": 0.0081, "learning_rate": 3.019779227044398e-06, "epoch": 2.2636655948553055, "percentage": 46.32, "elapsed_time": "0:16:19", "remaining_time": "0:18:55", "throughput": "615.45", "total_tokens": 602672}
{"current_steps": 89, "total_steps": 190, "loss": 0.0149, "learning_rate": 2.9770224884413625e-06, "epoch": 2.289389067524116, "percentage": 46.84, "elapsed_time": "0:16:30", "remaining_time": "0:18:43", "throughput": "615.35", "total_tokens": 609424}
{"current_steps": 90, "total_steps": 190, "loss": 0.001, "learning_rate": 2.9341204441673267e-06, "epoch": 2.315112540192926, "percentage": 47.37, "elapsed_time": "0:16:41", "remaining_time": "0:18:32", "throughput": "615.53", "total_tokens": 616448}
{"current_steps": 91, "total_steps": 190, "loss": 0.007, "learning_rate": 2.8910861626005774e-06, "epoch": 2.3408360128617365, "percentage": 47.89, "elapsed_time": "0:16:52", "remaining_time": "0:18:21", "throughput": "615.54", "total_tokens": 623296}
{"current_steps": 92, "total_steps": 190, "loss": 0.0089, "learning_rate": 2.847932752400164e-06, "epoch": 2.3665594855305465, "percentage": 48.42, "elapsed_time": "0:17:03", "remaining_time": "0:18:10", "throughput": "615.48", "total_tokens": 630064}
{"current_steps": 93, "total_steps": 190, "loss": 0.0013, "learning_rate": 2.804673358512869e-06, "epoch": 2.392282958199357, "percentage": 48.95, "elapsed_time": "0:17:14", "remaining_time": "0:17:59", "throughput": "615.67", "total_tokens": 637088}
{"current_steps": 94, "total_steps": 190, "loss": 0.0267, "learning_rate": 2.761321158169134e-06, "epoch": 2.418006430868167, "percentage": 49.47, "elapsed_time": "0:17:25", "remaining_time": "0:17:48", "throughput": "615.82", "total_tokens": 644080}
{"current_steps": 95, "total_steps": 190, "loss": 0.0171, "learning_rate": 2.717889356869146e-06, "epoch": 2.4437299035369775, "percentage": 50.0, "elapsed_time": "0:17:36", "remaining_time": "0:17:36", "throughput": "615.76", "total_tokens": 650848}
{"current_steps": 96, "total_steps": 190, "loss": 0.0375, "learning_rate": 2.6743911843603134e-06, "epoch": 2.469453376205788, "percentage": 50.53, "elapsed_time": "0:17:48", "remaining_time": "0:17:25", "throughput": "615.50", "total_tokens": 657424}
{"current_steps": 97, "total_steps": 190, "loss": 0.0101, "learning_rate": 2.6308398906073603e-06, "epoch": 2.495176848874598, "percentage": 51.05, "elapsed_time": "0:17:59", "remaining_time": "0:17:14", "throughput": "615.37", "total_tokens": 664128}
{"current_steps": 98, "total_steps": 190, "loss": 0.0282, "learning_rate": 2.587248741756253e-06, "epoch": 2.5209003215434085, "percentage": 51.58, "elapsed_time": "0:18:10", "remaining_time": "0:17:03", "throughput": "615.50", "total_tokens": 671120}
{"current_steps": 99, "total_steps": 190, "loss": 0.0069, "learning_rate": 2.543631016093209e-06, "epoch": 2.5466237942122185, "percentage": 52.11, "elapsed_time": "0:18:21", "remaining_time": "0:16:52", "throughput": "615.47", "total_tokens": 677920}
{"current_steps": 100, "total_steps": 190, "loss": 0.0135, "learning_rate": 2.5e-06, "epoch": 2.572347266881029, "percentage": 52.63, "elapsed_time": "0:18:32", "remaining_time": "0:16:41", "throughput": "615.66", "total_tokens": 684960}
{"current_steps": 101, "total_steps": 190, "loss": 0.0062, "learning_rate": 2.4563689839067913e-06, "epoch": 2.598070739549839, "percentage": 53.16, "elapsed_time": "0:18:43", "remaining_time": "0:16:30", "throughput": "615.71", "total_tokens": 691856}
{"current_steps": 102, "total_steps": 190, "loss": 0.005, "learning_rate": 2.4127512582437486e-06, "epoch": 2.6237942122186495, "percentage": 53.68, "elapsed_time": "0:18:54", "remaining_time": "0:16:19", "throughput": "615.56", "total_tokens": 698512}
{"current_steps": 103, "total_steps": 190, "loss": 0.0285, "learning_rate": 2.3691601093926406e-06, "epoch": 2.64951768488746, "percentage": 54.21, "elapsed_time": "0:19:05", "remaining_time": "0:16:07", "throughput": "615.65", "total_tokens": 705440}
{"current_steps": 104, "total_steps": 190, "loss": 0.0225, "learning_rate": 2.325608815639687e-06, "epoch": 2.67524115755627, "percentage": 54.74, "elapsed_time": "0:19:16", "remaining_time": "0:15:56", "throughput": "615.86", "total_tokens": 712528}
{"current_steps": 105, "total_steps": 190, "loss": 0.028, "learning_rate": 2.2821106431308546e-06, "epoch": 2.7009646302250805, "percentage": 55.26, "elapsed_time": "0:19:28", "remaining_time": "0:15:45", "throughput": "615.69", "total_tokens": 719168}
{"current_steps": 106, "total_steps": 190, "loss": 0.0176, "learning_rate": 2.238678841830867e-06, "epoch": 2.7266881028938905, "percentage": 55.79, "elapsed_time": "0:19:39", "remaining_time": "0:15:34", "throughput": "615.60", "total_tokens": 725904}
{"current_steps": 107, "total_steps": 190, "loss": 0.0047, "learning_rate": 2.195326641487132e-06, "epoch": 2.752411575562701, "percentage": 56.32, "elapsed_time": "0:19:50", "remaining_time": "0:15:23", "throughput": "615.36", "total_tokens": 732480}
{"current_steps": 108, "total_steps": 190, "loss": 0.0135, "learning_rate": 2.1520672475998374e-06, "epoch": 2.778135048231511, "percentage": 56.84, "elapsed_time": "0:20:01", "remaining_time": "0:15:12", "throughput": "615.25", "total_tokens": 739184}
{"current_steps": 109, "total_steps": 190, "loss": 0.0044, "learning_rate": 2.1089138373994226e-06, "epoch": 2.8038585209003215, "percentage": 57.37, "elapsed_time": "0:20:12", "remaining_time": "0:15:01", "throughput": "615.51", "total_tokens": 746320}
{"current_steps": 110, "total_steps": 190, "loss": 0.0252, "learning_rate": 2.0658795558326745e-06, "epoch": 2.829581993569132, "percentage": 57.89, "elapsed_time": "0:20:23", "remaining_time": "0:14:49", "throughput": "615.50", "total_tokens": 753136}
{"current_steps": 111, "total_steps": 190, "loss": 0.0249, "learning_rate": 2.022977511558638e-06, "epoch": 2.855305466237942, "percentage": 58.42, "elapsed_time": "0:20:34", "remaining_time": "0:14:38", "throughput": "615.61", "total_tokens": 760096}
{"current_steps": 112, "total_steps": 190, "loss": 0.0146, "learning_rate": 1.9802207729556023e-06, "epoch": 2.8810289389067525, "percentage": 58.95, "elapsed_time": "0:20:45", "remaining_time": "0:14:27", "throughput": "615.75", "total_tokens": 767104}
{"current_steps": 113, "total_steps": 190, "loss": 0.0044, "learning_rate": 1.937622364140338e-06, "epoch": 2.906752411575563, "percentage": 59.47, "elapsed_time": "0:20:56", "remaining_time": "0:14:16", "throughput": "615.69", "total_tokens": 773872}
{"current_steps": 114, "total_steps": 190, "loss": 0.0054, "learning_rate": 1.895195261000831e-06, "epoch": 2.932475884244373, "percentage": 60.0, "elapsed_time": "0:21:08", "remaining_time": "0:14:05", "throughput": "615.71", "total_tokens": 780736}
{"current_steps": 115, "total_steps": 190, "loss": 0.0106, "learning_rate": 1.852952387243698e-06, "epoch": 2.958199356913183, "percentage": 60.53, "elapsed_time": "0:21:19", "remaining_time": "0:13:54", "throughput": "615.58", "total_tokens": 787424}
{"current_steps": 116, "total_steps": 190, "loss": 0.0167, "learning_rate": 1.8109066104575023e-06, "epoch": 2.9839228295819935, "percentage": 61.05, "elapsed_time": "0:21:30", "remaining_time": "0:13:43", "throughput": "615.53", "total_tokens": 794224}
{"current_steps": 117, "total_steps": 190, "loss": 0.009, "learning_rate": 1.7690707381931585e-06, "epoch": 3.009646302250804, "percentage": 61.58, "elapsed_time": "0:21:41", "remaining_time": "0:13:31", "throughput": "615.55", "total_tokens": 801088}
{"current_steps": 118, "total_steps": 190, "loss": 0.0024, "learning_rate": 1.7274575140626318e-06, "epoch": 3.035369774919614, "percentage": 62.11, "elapsed_time": "0:21:52", "remaining_time": "0:13:20", "throughput": "615.66", "total_tokens": 808048}
{"current_steps": 119, "total_steps": 190, "loss": 0.0235, "learning_rate": 1.686079613857109e-06, "epoch": 3.0610932475884245, "percentage": 62.63, "elapsed_time": "0:22:03", "remaining_time": "0:13:09", "throughput": "615.60", "total_tokens": 814800}
{"current_steps": 120, "total_steps": 190, "loss": 0.0179, "learning_rate": 1.6449496416858285e-06, "epoch": 3.0868167202572345, "percentage": 63.16, "elapsed_time": "0:22:14", "remaining_time": "0:12:58", "throughput": "615.53", "total_tokens": 821536}
{"current_steps": 121, "total_steps": 190, "loss": 0.0059, "learning_rate": 1.6040801261367494e-06, "epoch": 3.112540192926045, "percentage": 63.68, "elapsed_time": "0:22:25", "remaining_time": "0:12:47", "throughput": "615.35", "total_tokens": 828128}
{"current_steps": 122, "total_steps": 190, "loss": 0.0017, "learning_rate": 1.56348351646022e-06, "epoch": 3.1382636655948555, "percentage": 64.21, "elapsed_time": "0:22:36", "remaining_time": "0:12:36", "throughput": "615.08", "total_tokens": 834608}
{"current_steps": 123, "total_steps": 190, "loss": 0.0018, "learning_rate": 1.5231721787768162e-06, "epoch": 3.1639871382636655, "percentage": 64.74, "elapsed_time": "0:22:48", "remaining_time": "0:12:25", "throughput": "615.02", "total_tokens": 841360}
{"current_steps": 124, "total_steps": 190, "loss": 0.0032, "learning_rate": 1.4831583923105e-06, "epoch": 3.189710610932476, "percentage": 65.26, "elapsed_time": "0:22:59", "remaining_time": "0:12:14", "throughput": "615.23", "total_tokens": 848496}
{"current_steps": 125, "total_steps": 190, "loss": 0.0019, "learning_rate": 1.443454345648252e-06, "epoch": 3.215434083601286, "percentage": 65.79, "elapsed_time": "0:23:10", "remaining_time": "0:12:02", "throughput": "615.29", "total_tokens": 855424}
{"current_steps": 126, "total_steps": 190, "loss": 0.0014, "learning_rate": 1.4040721330273063e-06, "epoch": 3.2411575562700965, "percentage": 66.32, "elapsed_time": "0:23:21", "remaining_time": "0:11:51", "throughput": "615.27", "total_tokens": 862224}
{"current_steps": 127, "total_steps": 190, "loss": 0.0052, "learning_rate": 1.3650237506511333e-06, "epoch": 3.266881028938907, "percentage": 66.84, "elapsed_time": "0:23:32", "remaining_time": "0:11:40", "throughput": "615.45", "total_tokens": 869312}
{"current_steps": 128, "total_steps": 190, "loss": 0.0005, "learning_rate": 1.3263210930352737e-06, "epoch": 3.292604501607717, "percentage": 67.37, "elapsed_time": "0:23:43", "remaining_time": "0:11:29", "throughput": "615.55", "total_tokens": 876272}
{"current_steps": 129, "total_steps": 190, "loss": 0.0131, "learning_rate": 1.2879759493841577e-06, "epoch": 3.3183279742765275, "percentage": 67.89, "elapsed_time": "0:23:54", "remaining_time": "0:11:18", "throughput": "615.52", "total_tokens": 883072}
{"current_steps": 130, "total_steps": 190, "loss": 0.0009, "learning_rate": 1.2500000000000007e-06, "epoch": 3.3440514469453375, "percentage": 68.42, "elapsed_time": "0:24:05", "remaining_time": "0:11:07", "throughput": "615.54", "total_tokens": 889920}
{"current_steps": 131, "total_steps": 190, "loss": 0.0057, "learning_rate": 1.2124048127248644e-06, "epoch": 3.369774919614148, "percentage": 68.95, "elapsed_time": "0:24:16", "remaining_time": "0:10:56", "throughput": "615.63", "total_tokens": 896896}
{"current_steps": 132, "total_steps": 190, "loss": 0.0002, "learning_rate": 1.1752018394169882e-06, "epoch": 3.395498392282958, "percentage": 69.47, "elapsed_time": "0:24:27", "remaining_time": "0:10:45", "throughput": "615.53", "total_tokens": 903600}
{"current_steps": 133, "total_steps": 190, "loss": 0.0002, "learning_rate": 1.1384024124624324e-06, "epoch": 3.4212218649517685, "percentage": 70.0, "elapsed_time": "0:24:39", "remaining_time": "0:10:33", "throughput": "615.37", "total_tokens": 910208}
{"current_steps": 134, "total_steps": 190, "loss": 0.0145, "learning_rate": 1.1020177413231334e-06, "epoch": 3.446945337620579, "percentage": 70.53, "elapsed_time": "0:24:50", "remaining_time": "0:10:22", "throughput": "615.56", "total_tokens": 917328}
{"current_steps": 135, "total_steps": 190, "loss": 0.0034, "learning_rate": 1.0660589091223854e-06, "epoch": 3.472668810289389, "percentage": 71.05, "elapsed_time": "0:25:01", "remaining_time": "0:10:11", "throughput": "615.59", "total_tokens": 924192}
{"current_steps": 136, "total_steps": 190, "loss": 0.0156, "learning_rate": 1.0305368692688175e-06, "epoch": 3.4983922829581995, "percentage": 71.58, "elapsed_time": "0:25:12", "remaining_time": "0:10:00", "throughput": "615.43", "total_tokens": 930784}
{"current_steps": 137, "total_steps": 190, "loss": 0.0013, "learning_rate": 9.95462442119879e-07, "epoch": 3.5241157556270095, "percentage": 72.11, "elapsed_time": "0:25:23", "remaining_time": "0:09:49", "throughput": "615.59", "total_tokens": 937856}
{"current_steps": 138, "total_steps": 190, "loss": 0.0007, "learning_rate": 9.608463116858544e-07, "epoch": 3.54983922829582, "percentage": 72.63, "elapsed_time": "0:25:34", "remaining_time": "0:09:38", "throughput": "615.56", "total_tokens": 944640}
{"current_steps": 139, "total_steps": 190, "loss": 0.0005, "learning_rate": 9.266990223754069e-07, "epoch": 3.57556270096463, "percentage": 73.16, "elapsed_time": "0:25:45", "remaining_time": "0:09:27", "throughput": "615.58", "total_tokens": 951504}
{"current_steps": 140, "total_steps": 190, "loss": 0.0034, "learning_rate": 8.930309757836517e-07, "epoch": 3.6012861736334405, "percentage": 73.68, "elapsed_time": "0:25:56", "remaining_time": "0:09:16", "throughput": "615.52", "total_tokens": 958240}
{"current_steps": 141, "total_steps": 190, "loss": 0.0001, "learning_rate": 8.598524275237321e-07, "epoch": 3.627009646302251, "percentage": 74.21, "elapsed_time": "0:26:07", "remaining_time": "0:09:04", "throughput": "615.41", "total_tokens": 964912}
{"current_steps": 142, "total_steps": 190, "loss": 0.001, "learning_rate": 8.271734841028553e-07, "epoch": 3.652733118971061, "percentage": 74.74, "elapsed_time": "0:26:19", "remaining_time": "0:08:53", "throughput": "615.49", "total_tokens": 971872}
{"current_steps": 143, "total_steps": 190, "loss": 0.0123, "learning_rate": 7.950040998437541e-07, "epoch": 3.6784565916398715, "percentage": 75.26, "elapsed_time": "0:26:30", "remaining_time": "0:08:42", "throughput": "615.44", "total_tokens": 978640}
{"current_steps": 144, "total_steps": 190, "loss": 0.0002, "learning_rate": 7.633540738525066e-07, "epoch": 3.7041800643086815, "percentage": 75.79, "elapsed_time": "0:26:41", "remaining_time": "0:08:31", "throughput": "615.35", "total_tokens": 985328}
{"current_steps": 145, "total_steps": 190, "loss": 0.011, "learning_rate": 7.322330470336314e-07, "epoch": 3.729903536977492, "percentage": 76.32, "elapsed_time": "0:26:52", "remaining_time": "0:08:20", "throughput": "615.39", "total_tokens": 992224}
{"current_steps": 146, "total_steps": 190, "loss": 0.0008, "learning_rate": 7.016504991533727e-07, "epoch": 3.755627009646302, "percentage": 76.84, "elapsed_time": "0:27:03", "remaining_time": "0:08:09", "throughput": "615.17", "total_tokens": 998688}
{"current_steps": 147, "total_steps": 190, "loss": 0.0003, "learning_rate": 6.716157459520739e-07, "epoch": 3.7813504823151125, "percentage": 77.37, "elapsed_time": "0:27:14", "remaining_time": "0:07:58", "throughput": "615.50", "total_tokens": 1006032}
{"current_steps": 148, "total_steps": 190, "loss": 0.0018, "learning_rate": 6.421379363065142e-07, "epoch": 3.807073954983923, "percentage": 77.89, "elapsed_time": "0:27:25", "remaining_time": "0:07:46", "throughput": "615.55", "total_tokens": 1012944}
{"current_steps": 149, "total_steps": 190, "loss": 0.0016, "learning_rate": 6.1322604944307e-07, "epoch": 3.832797427652733, "percentage": 78.42, "elapsed_time": "0:27:36", "remaining_time": "0:07:35", "throughput": "615.59", "total_tokens": 1019856}
{"current_steps": 150, "total_steps": 190, "loss": 0.0021, "learning_rate": 5.848888922025553e-07, "epoch": 3.8585209003215435, "percentage": 78.95, "elapsed_time": "0:27:47", "remaining_time": "0:07:24", "throughput": "615.56", "total_tokens": 1026656}
{"current_steps": 151, "total_steps": 190, "loss": 0.0001, "learning_rate": 5.571350963575728e-07, "epoch": 3.884244372990354, "percentage": 79.47, "elapsed_time": "0:27:58", "remaining_time": "0:07:13", "throughput": "615.68", "total_tokens": 1033696}
{"current_steps": 152, "total_steps": 190, "loss": 0.0001, "learning_rate": 5.299731159831953e-07, "epoch": 3.909967845659164, "percentage": 80.0, "elapsed_time": "0:28:10", "remaining_time": "0:07:02", "throughput": "615.50", "total_tokens": 1040240}
{"current_steps": 153, "total_steps": 190, "loss": 0.0003, "learning_rate": 5.034112248817685e-07, "epoch": 3.935691318327974, "percentage": 80.53, "elapsed_time": "0:28:21", "remaining_time": "0:06:51", "throughput": "615.45", "total_tokens": 1046992}
{"current_steps": 154, "total_steps": 190, "loss": 0.0002, "learning_rate": 4.774575140626317e-07, "epoch": 3.9614147909967845, "percentage": 81.05, "elapsed_time": "0:28:32", "remaining_time": "0:06:40", "throughput": "615.52", "total_tokens": 1053936}
{"current_steps": 155, "total_steps": 190, "loss": 0.0002, "learning_rate": 4.5211988927752026e-07, "epoch": 3.987138263665595, "percentage": 81.58, "elapsed_time": "0:28:43", "remaining_time": "0:06:29", "throughput": "615.41", "total_tokens": 1060576}
{"current_steps": 156, "total_steps": 190, "loss": 0.0, "learning_rate": 4.27406068612396e-07, "epoch": 4.012861736334405, "percentage": 82.11, "elapsed_time": "0:28:54", "remaining_time": "0:06:18", "throughput": "615.56", "total_tokens": 1067648}
{"current_steps": 157, "total_steps": 190, "loss": 0.0003, "learning_rate": 4.033235801364402e-07, "epoch": 4.038585209003215, "percentage": 82.63, "elapsed_time": "0:29:05", "remaining_time": "0:06:06", "throughput": "615.58", "total_tokens": 1074512}
{"current_steps": 158, "total_steps": 190, "loss": 0.0002, "learning_rate": 3.798797596089351e-07, "epoch": 4.064308681672026, "percentage": 83.16, "elapsed_time": "0:29:16", "remaining_time": "0:05:55", "throughput": "615.55", "total_tokens": 1081296}
{"current_steps": 159, "total_steps": 190, "loss": 0.0001, "learning_rate": 3.5708174824471947e-07, "epoch": 4.090032154340836, "percentage": 83.68, "elapsed_time": "0:29:27", "remaining_time": "0:05:44", "throughput": "615.40", "total_tokens": 1087888}
{"current_steps": 160, "total_steps": 190, "loss": 0.0001, "learning_rate": 3.3493649053890325e-07, "epoch": 4.115755627009646, "percentage": 84.21, "elapsed_time": "0:29:38", "remaining_time": "0:05:33", "throughput": "615.53", "total_tokens": 1094960}
{"current_steps": 161, "total_steps": 190, "loss": 0.0001, "learning_rate": 3.134507321515107e-07, "epoch": 4.141479099678457, "percentage": 84.74, "elapsed_time": "0:29:50", "remaining_time": "0:05:22", "throughput": "615.52", "total_tokens": 1101776}
{"current_steps": 162, "total_steps": 190, "loss": 0.0, "learning_rate": 2.9263101785268253e-07, "epoch": 4.167202572347267, "percentage": 85.26, "elapsed_time": "0:30:01", "remaining_time": "0:05:11", "throughput": "615.59", "total_tokens": 1108736}
{"current_steps": 163, "total_steps": 190, "loss": 0.0001, "learning_rate": 2.7248368952908055e-07, "epoch": 4.192926045016077, "percentage": 85.79, "elapsed_time": "0:30:12", "remaining_time": "0:05:00", "throughput": "615.69", "total_tokens": 1115744}
{"current_steps": 164, "total_steps": 190, "loss": 0.0, "learning_rate": 2.53014884252083e-07, "epoch": 4.218649517684887, "percentage": 86.32, "elapsed_time": "0:30:23", "remaining_time": "0:04:49", "throughput": "615.71", "total_tokens": 1122592}
{"current_steps": 165, "total_steps": 190, "loss": 0.0001, "learning_rate": 2.3423053240837518e-07, "epoch": 4.244372990353698, "percentage": 86.84, "elapsed_time": "0:30:34", "remaining_time": "0:04:37", "throughput": "615.55", "total_tokens": 1129136}
{"current_steps": 166, "total_steps": 190, "loss": 0.0001, "learning_rate": 2.1613635589349756e-07, "epoch": 4.270096463022508, "percentage": 87.37, "elapsed_time": "0:30:45", "remaining_time": "0:04:26", "throughput": "615.61", "total_tokens": 1136080}
{"current_steps": 167, "total_steps": 190, "loss": 0.0002, "learning_rate": 1.9873786636889908e-07, "epoch": 4.295819935691318, "percentage": 87.89, "elapsed_time": "0:30:56", "remaining_time": "0:04:15", "throughput": "615.64", "total_tokens": 1142976}
{"current_steps": 168, "total_steps": 190, "loss": 0.0002, "learning_rate": 1.8204036358303173e-07, "epoch": 4.321543408360129, "percentage": 88.42, "elapsed_time": "0:31:07", "remaining_time": "0:04:04", "throughput": "615.58", "total_tokens": 1149712}
{"current_steps": 169, "total_steps": 190, "loss": 0.0001, "learning_rate": 1.6604893375699594e-07, "epoch": 4.347266881028939, "percentage": 88.95, "elapsed_time": "0:31:18", "remaining_time": "0:03:53", "throughput": "615.46", "total_tokens": 1156336}
{"current_steps": 170, "total_steps": 190, "loss": 0.0001, "learning_rate": 1.507684480352292e-07, "epoch": 4.372990353697749, "percentage": 89.47, "elapsed_time": "0:31:29", "remaining_time": "0:03:42", "throughput": "615.43", "total_tokens": 1163120}
{"current_steps": 171, "total_steps": 190, "loss": 0.0115, "learning_rate": 1.362035610017079e-07, "epoch": 4.39871382636656, "percentage": 90.0, "elapsed_time": "0:31:41", "remaining_time": "0:03:31", "throughput": "615.48", "total_tokens": 1170032}
{"current_steps": 172, "total_steps": 190, "loss": 0.0005, "learning_rate": 1.223587092621162e-07, "epoch": 4.42443729903537, "percentage": 90.53, "elapsed_time": "0:31:52", "remaining_time": "0:03:20", "throughput": "615.47", "total_tokens": 1176832}
{"current_steps": 173, "total_steps": 190, "loss": 0.0003, "learning_rate": 1.0923811009241142e-07, "epoch": 4.45016077170418, "percentage": 91.05, "elapsed_time": "0:32:03", "remaining_time": "0:03:08", "throughput": "615.57", "total_tokens": 1183856}
{"current_steps": 174, "total_steps": 190, "loss": 0.005, "learning_rate": 9.684576015420277e-08, "epoch": 4.47588424437299, "percentage": 91.58, "elapsed_time": "0:32:14", "remaining_time": "0:02:57", "throughput": "615.49", "total_tokens": 1190544}
{"current_steps": 175, "total_steps": 190, "loss": 0.0003, "learning_rate": 8.518543427732951e-08, "epoch": 4.501607717041801, "percentage": 92.11, "elapsed_time": "0:32:25", "remaining_time": "0:02:46", "throughput": "615.38", "total_tokens": 1197168}
{"current_steps": 176, "total_steps": 190, "loss": 0.0008, "learning_rate": 7.426068431000883e-08, "epoch": 4.527331189710611, "percentage": 92.63, "elapsed_time": "0:32:36", "remaining_time": "0:02:35", "throughput": "615.27", "total_tokens": 1203776}
{"current_steps": 177, "total_steps": 190, "loss": 0.0, "learning_rate": 6.407483803691216e-08, "epoch": 4.553054662379421, "percentage": 93.16, "elapsed_time": "0:32:47", "remaining_time": "0:02:24", "throughput": "615.37", "total_tokens": 1210816}
{"current_steps": 178, "total_steps": 190, "loss": 0.0, "learning_rate": 5.463099816548578e-08, "epoch": 4.578778135048232, "percentage": 93.68, "elapsed_time": "0:32:58", "remaining_time": "0:02:13", "throughput": "615.39", "total_tokens": 1217696}
{"current_steps": 179, "total_steps": 190, "loss": 0.0042, "learning_rate": 4.593204138084006e-08, "epoch": 4.604501607717042, "percentage": 94.21, "elapsed_time": "0:33:09", "remaining_time": "0:02:02", "throughput": "615.48", "total_tokens": 1224704}
{"current_steps": 180, "total_steps": 190, "loss": 0.0004, "learning_rate": 3.798061746947995e-08, "epoch": 4.630225080385852, "percentage": 94.74, "elapsed_time": "0:33:20", "remaining_time": "0:01:51", "throughput": "615.54", "total_tokens": 1231664}
{"current_steps": 181, "total_steps": 190, "loss": 0.0001, "learning_rate": 3.077914851215585e-08, "epoch": 4.655948553054662, "percentage": 95.26, "elapsed_time": "0:33:32", "remaining_time": "0:01:40", "throughput": "615.42", "total_tokens": 1238240}
{"current_steps": 182, "total_steps": 190, "loss": 0.0001, "learning_rate": 2.4329828146074096e-08, "epoch": 4.681672025723473, "percentage": 95.79, "elapsed_time": "0:33:43", "remaining_time": "0:01:28", "throughput": "615.30", "total_tokens": 1244832}
{"current_steps": 183, "total_steps": 190, "loss": 0.0004, "learning_rate": 1.8634620896695044e-08, "epoch": 4.707395498392283, "percentage": 96.32, "elapsed_time": "0:33:54", "remaining_time": "0:01:17", "throughput": "615.12", "total_tokens": 1251296}
{"current_steps": 184, "total_steps": 190, "loss": 0.0, "learning_rate": 1.3695261579316776e-08, "epoch": 4.733118971061093, "percentage": 96.84, "elapsed_time": "0:34:05", "remaining_time": "0:01:06", "throughput": "615.05", "total_tokens": 1257968}
{"current_steps": 185, "total_steps": 190, "loss": 0.0009, "learning_rate": 9.513254770636138e-09, "epoch": 4.758842443729904, "percentage": 97.37, "elapsed_time": "0:34:16", "remaining_time": "0:00:55", "throughput": "615.11", "total_tokens": 1264928}
{"current_steps": 186, "total_steps": 190, "loss": 0.0001, "learning_rate": 6.089874350439507e-09, "epoch": 4.784565916398714, "percentage": 97.89, "elapsed_time": "0:34:27", "remaining_time": "0:00:44", "throughput": "615.12", "total_tokens": 1271792}
{"current_steps": 187, "total_steps": 190, "loss": 0.0004, "learning_rate": 3.4261631135654174e-09, "epoch": 4.810289389067524, "percentage": 98.42, "elapsed_time": "0:34:38", "remaining_time": "0:00:33", "throughput": "615.30", "total_tokens": 1279024}
{"current_steps": 188, "total_steps": 190, "loss": 0.0002, "learning_rate": 1.5229324522605949e-09, "epoch": 4.836012861736334, "percentage": 98.95, "elapsed_time": "0:34:49", "remaining_time": "0:00:22", "throughput": "615.28", "total_tokens": 1285792}
{"current_steps": 189, "total_steps": 190, "loss": 0.0, "learning_rate": 3.8076210902182607e-10, "epoch": 4.861736334405145, "percentage": 99.47, "elapsed_time": "0:35:00", "remaining_time": "0:00:11", "throughput": "615.26", "total_tokens": 1292576}
{"current_steps": 190, "total_steps": 190, "loss": 0.0001, "learning_rate": 0.0, "epoch": 4.887459807073955, "percentage": 100.0, "elapsed_time": "0:35:11", "remaining_time": "0:00:00", "throughput": "615.25", "total_tokens": 1299392}
{"current_steps": 190, "total_steps": 190, "epoch": 4.887459807073955, "percentage": 100.0, "elapsed_time": "0:36:02", "remaining_time": "0:00:00", "throughput": "600.99", "total_tokens": 1299392}