Training in progress, epoch 1

Files changed (9) hide show

merges.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

pytorch_model.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:8c8a122f3f4f86d688c0bab2b58ac19ccc6a658862730fbdb6adb5a38a9b65b6
 size 510413893

 version https://git-lfs.github.com/spec/v1
+oid sha256:cbb1c30cdf34202911ba86f2bd31a42ff658705416f1bd9a10c5bc9cb42db4c4
 size 510413893

runs/May06_12-52-26_0f750b5bcb8c/1683377843.9893055/events.out.tfevents.1683377843.0f750b5bcb8c.194.1 ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:c7bb3c0e47566184749c104cb4e5b370eaeb2618693681e7d5f5cb7e3ea30b25
+size 5906

runs/May06_12-52-26_0f750b5bcb8c/events.out.tfevents.1683377843.0f750b5bcb8c.194.0 ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:60456f39cc4fc49fc827bc6b96e379cc1181cadb2769437dfcbaa62ba78b0ac6
+size 5593

special_tokens_map.json CHANGED Viewed

@@ -1,15 +1,6 @@
 {
-  "bos_token": "[CLS]",
-  "cls_token": "[CLS]",
-  "eos_token": "[SEP]",
-  "mask_token": {
-    "content": "[MASK]",
-    "lstrip": true,
-    "normalized": false,
-    "rstrip": false,
-    "single_word": false
-  },
-  "pad_token": "<pad>",
-  "sep_token": "[SEP]",
-  "unk_token": "<unk>"
 }

 {
+  "bos_token": "<|endoftext|>",
+  "eos_token": "<|endoftext|>",
+  "pad_token": "<|endoftext|>",
+  "unk_token": "<|endoftext|>"
 }

tokenizer.json CHANGED Viewed

The diff for this file is too large to render. See raw diff

tokenizer_config.json CHANGED Viewed

@@ -1,23 +1,9 @@
 {
   "add_prefix_space": true,
-  "bos_token": "[CLS]",
   "clean_up_tokenization_spaces": true,
-  "cls_token": "[CLS]",
-  "do_lower_case": true,
-  "eos_token": "[SEP]",
-  "keep_accents": false,
-  "mask_token": {
-    "__type": "AddedToken",
-    "content": "[MASK]",
-    "lstrip": true,
-    "normalized": false,
-    "rstrip": false,
-    "single_word": false
-  },
-  "model_max_length": 512,
-  "pad_token": "<pad>",
-  "remove_space": true,
-  "sep_token": "[SEP]",
-  "tokenizer_class": "AlbertTokenizer",
-  "unk_token": "<unk>"
 }

 {
   "add_prefix_space": true,
+  "bos_token": "<|endoftext|>",
   "clean_up_tokenization_spaces": true,
+  "eos_token": "<|endoftext|>",
+  "model_max_length": 1024,
+  "tokenizer_class": "GPT2Tokenizer",
+  "unk_token": "<|endoftext|>"
 }

training_args.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:aa89d6a8d99d8406db39235d20bc85b05d2b79af5f16be28371ef654be5fbbe8
 size 3579

 version https://git-lfs.github.com/spec/v1
+oid sha256:0b401c157dc8f2213d8f8acf2729efc3d181d90cc16999e83b8103f57eb391be
 size 3579

vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff