mpnet_pretrain_from_16 / tokenizer_config.json
onlydj96's picture
add tokenizer
82e845b
raw
history blame contribute delete
No virus
1.59 kB
{
"bos_token": {
"__type": "AddedToken",
"content": "<s>",
"lstrip": false,
"normalized": true,
"rstrip": false,
"single_word": false
},
"cls_token": {
"__type": "AddedToken",
"content": "[CLS]",
"lstrip": false,
"normalized": true,
"rstrip": false,
"single_word": false
},
"do_basic_tokenize": true,
"do_lower_case": false,
"eos_token": {
"__type": "AddedToken",
"content": "</s>",
"lstrip": false,
"normalized": true,
"rstrip": false,
"single_word": false
},
"mask_token": {
"__type": "AddedToken",
"content": "[MASK]",
"lstrip": true,
"normalized": true,
"rstrip": false,
"single_word": false
},
"name_or_path": "onlydj96/mpnet_pretrain_first_train",
"never_split": null,
"pad_token": {
"__type": "AddedToken",
"content": "[PAD]",
"lstrip": false,
"normalized": true,
"rstrip": false,
"single_word": false
},
"sep_token": {
"__type": "AddedToken",
"content": "[SEP]",
"lstrip": false,
"normalized": true,
"rstrip": false,
"single_word": false
},
"special_tokens_map_file": "/root/.cache/huggingface/transformers/648a3d5f7e33d9edb597394430c91a1e7e685d84fae15f7cf40d86e408bb5bb3.1b83d0d7f4d455d37c683966d465a99be7f33983cf93b19ad8d2d23d044ea57a",
"strip_accents": null,
"tokenize_chinese_chars": true,
"tokenizer_class": "MPNetTokenizer",
"unk_token": {
"__type": "AddedToken",
"content": "[UNK]",
"lstrip": false,
"normalized": true,
"rstrip": false,
"single_word": false
}
}