avuhong commited on
Commit
46656c3
1 Parent(s): d04afdf

add tokenizer

Browse files
Files changed (3) hide show
  1. special_tokens_map.json +1 -0
  2. tokenizer_config.json +1 -0
  3. vocab.txt +33 -0
special_tokens_map.json ADDED
@@ -0,0 +1 @@
 
1
+ {"eos_token": "<eos>", "unk_token": "<unk>", "pad_token": "<pad>", "cls_token": "<cls>", "mask_token": "<mask>"}
tokenizer_config.json ADDED
@@ -0,0 +1 @@
 
1
+ {"do_lower_case": false, "special_tokens_map_file": "/root/.cache/huggingface/transformers/7b44bebd885ea087e4c615e7d33219653ea809b6aa3b3b5d9c1874b574f5905d.9c5e761c0365345e78df6372a132f3313e405fe3392827daf71113edd456660b", "tokenizer_file": null, "name_or_path": "facebook/esm-1b", "tokenizer_class": "ESMTokenizer"}
vocab.txt ADDED
@@ -0,0 +1,33 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+ <cls>
2
+ <pad>
3
+ <eos>
4
+ <unk>
5
+ L
6
+ A
7
+ G
8
+ V
9
+ S
10
+ E
11
+ R
12
+ T
13
+ I
14
+ D
15
+ P
16
+ K
17
+ Q
18
+ N
19
+ F
20
+ Y
21
+ M
22
+ H
23
+ W
24
+ C
25
+ X
26
+ B
27
+ U
28
+ Z
29
+ O
30
+ .
31
+ -
32
+ <null_1>
33
+ <mask>