hgharibi's picture
add tokenizer
ba0b20c
{"ا": 1, "ب": 2, "ت": 3, "ث": 4, "ج": 5, "ح": 6, "خ": 7, "د": 8, "ذ": 9, "ر": 10, "ز": 11, "س": 12, "ش": 13, "ص": 14, "ض": 15, "ط": 16, "ظ": 17, "ع": 18, "غ": 19, "ف": 20, "ق": 21, "ل": 22, "م": 23, "ن": 24, "ه": 25, "و": 26, "ى": 27, "ٓ": 28, "پ": 29, "چ": 30, "ژ": 31, "ک": 32, "گ": 33, "ی": 34, "|": 0, "[UNK]": 35, "[PAD]": 36}