add tokenizer
Browse files- vocab.json +1 -1
vocab.json
CHANGED
@@ -1 +1 @@
|
|
1 |
-
{"
|
|
|
1 |
+
{"ш": 0, "ж": 2, "с": 3, "д": 4, "я": 5, "и": 6, "ь": 7, "х": 8, "н": 9, "е": 10, "к": 11, "ф": 12, "ъ": 13, "б": 14, "у": 15, "ц": 16, "ю": 17, "м": 18, "п": 19, "й": 20, "ы": 21, "ё": 22, "в": 23, "о": 24, "з": 25, "а": 26, "р": 27, "т": 28, "щ": 29, "ч": 30, "л": 31, "г": 32, "э": 33, "|": 1, "[UNK]": 34, "[PAD]": 35}
|