jteng commited on
Commit
88a7e96
1 Parent(s): 108653f

Model save

Browse files
config.json CHANGED
@@ -1,30 +1,24 @@
1
  {
2
- "_name_or_path": "deepset/roberta-base-squad2",
 
3
  "architectures": [
4
- "RobertaForQuestionAnswering"
5
  ],
6
- "attention_probs_dropout_prob": 0.1,
7
- "bos_token_id": 0,
8
- "classifier_dropout": null,
9
- "eos_token_id": 2,
10
- "gradient_checkpointing": false,
11
- "hidden_act": "gelu",
12
- "hidden_dropout_prob": 0.1,
13
- "hidden_size": 768,
14
  "initializer_range": 0.02,
15
- "intermediate_size": 3072,
16
- "language": "english",
17
- "layer_norm_eps": 1e-05,
18
- "max_position_embeddings": 514,
19
- "model_type": "roberta",
20
- "name": "Roberta",
21
- "num_attention_heads": 12,
22
- "num_hidden_layers": 12,
23
- "pad_token_id": 1,
24
- "position_embedding_type": "absolute",
25
  "torch_dtype": "float32",
26
  "transformers_version": "4.22.1",
27
- "type_vocab_size": 1,
28
- "use_cache": true,
29
- "vocab_size": 50265
30
  }
 
1
  {
2
+ "_name_or_path": "distilbert-base-uncased-distilled-squad",
3
+ "activation": "gelu",
4
  "architectures": [
5
+ "DistilBertForQuestionAnswering"
6
  ],
7
+ "attention_dropout": 0.1,
8
+ "dim": 768,
9
+ "dropout": 0.1,
10
+ "hidden_dim": 3072,
 
 
 
 
11
  "initializer_range": 0.02,
12
+ "max_position_embeddings": 512,
13
+ "model_type": "distilbert",
14
+ "n_heads": 12,
15
+ "n_layers": 6,
16
+ "pad_token_id": 0,
17
+ "qa_dropout": 0.1,
18
+ "seq_classif_dropout": 0.2,
19
+ "sinusoidal_pos_embds": false,
20
+ "tie_weights_": true,
 
21
  "torch_dtype": "float32",
22
  "transformers_version": "4.22.1",
23
+ "vocab_size": 30522
 
 
24
  }
pytorch_model.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:23606f03b73ecafd38e7d81a32c4fe62fab2eb8f1a90716a4850236413df52a5
3
- size 496297393
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2de7d80e285e2490976c145bf6deb6df19b9f92f9f723bcdda8c6cf47eef32e2
3
+ size 265489909
special_tokens_map.json CHANGED
@@ -1,51 +1,7 @@
1
  {
2
- "bos_token": {
3
- "content": "<s>",
4
- "lstrip": false,
5
- "normalized": true,
6
- "rstrip": false,
7
- "single_word": false
8
- },
9
- "cls_token": {
10
- "content": "<s>",
11
- "lstrip": false,
12
- "normalized": true,
13
- "rstrip": false,
14
- "single_word": false
15
- },
16
- "eos_token": {
17
- "content": "</s>",
18
- "lstrip": false,
19
- "normalized": true,
20
- "rstrip": false,
21
- "single_word": false
22
- },
23
- "mask_token": {
24
- "content": "<mask>",
25
- "lstrip": true,
26
- "normalized": true,
27
- "rstrip": false,
28
- "single_word": false
29
- },
30
- "pad_token": {
31
- "content": "<pad>",
32
- "lstrip": false,
33
- "normalized": true,
34
- "rstrip": false,
35
- "single_word": false
36
- },
37
- "sep_token": {
38
- "content": "</s>",
39
- "lstrip": false,
40
- "normalized": true,
41
- "rstrip": false,
42
- "single_word": false
43
- },
44
- "unk_token": {
45
- "content": "<unk>",
46
- "lstrip": false,
47
- "normalized": true,
48
- "rstrip": false,
49
- "single_word": false
50
- }
51
  }
 
1
  {
2
+ "cls_token": "[CLS]",
3
+ "mask_token": "[MASK]",
4
+ "pad_token": "[PAD]",
5
+ "sep_token": "[SEP]",
6
+ "unk_token": "[UNK]"
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
7
  }
tokenizer.json CHANGED
The diff for this file is too large to render. See raw diff
 
tokenizer_config.json CHANGED
@@ -1,67 +1,14 @@
1
  {
2
- "add_prefix_space": false,
3
- "bos_token": {
4
- "__type": "AddedToken",
5
- "content": "<s>",
6
- "lstrip": false,
7
- "normalized": true,
8
- "rstrip": false,
9
- "single_word": false
10
- },
11
- "cls_token": {
12
- "__type": "AddedToken",
13
- "content": "<s>",
14
- "lstrip": false,
15
- "normalized": true,
16
- "rstrip": false,
17
- "single_word": false
18
- },
19
- "do_lower_case": false,
20
- "eos_token": {
21
- "__type": "AddedToken",
22
- "content": "</s>",
23
- "lstrip": false,
24
- "normalized": true,
25
- "rstrip": false,
26
- "single_word": false
27
- },
28
- "errors": "replace",
29
- "full_tokenizer_file": null,
30
- "mask_token": {
31
- "__type": "AddedToken",
32
- "content": "<mask>",
33
- "lstrip": true,
34
- "normalized": true,
35
- "rstrip": false,
36
- "single_word": false
37
- },
38
  "model_max_length": 512,
39
- "name_or_path": "deepset/roberta-base-squad2",
40
- "pad_token": {
41
- "__type": "AddedToken",
42
- "content": "<pad>",
43
- "lstrip": false,
44
- "normalized": true,
45
- "rstrip": false,
46
- "single_word": false
47
- },
48
- "sep_token": {
49
- "__type": "AddedToken",
50
- "content": "</s>",
51
- "lstrip": false,
52
- "normalized": true,
53
- "rstrip": false,
54
- "single_word": false
55
- },
56
- "special_tokens_map_file": "/home/zxt180005/.cache/huggingface/hub/models--deepset--roberta-base-squad2/snapshots/65f7840c86b02ca4df86024defd6de99fcf1fc10/special_tokens_map.json",
57
- "tokenizer_class": "RobertaTokenizer",
58
- "trim_offsets": true,
59
- "unk_token": {
60
- "__type": "AddedToken",
61
- "content": "<unk>",
62
- "lstrip": false,
63
- "normalized": true,
64
- "rstrip": false,
65
- "single_word": false
66
- }
67
  }
 
1
  {
2
+ "cls_token": "[CLS]",
3
+ "do_lower_case": true,
4
+ "mask_token": "[MASK]",
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
5
  "model_max_length": 512,
6
+ "name_or_path": "distilbert-base-uncased-distilled-squad",
7
+ "pad_token": "[PAD]",
8
+ "sep_token": "[SEP]",
9
+ "special_tokens_map_file": null,
10
+ "strip_accents": null,
11
+ "tokenize_chinese_chars": true,
12
+ "tokenizer_class": "DistilBertTokenizer",
13
+ "unk_token": "[UNK]"
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
14
  }
training_args.bin CHANGED
@@ -1,3 +1,3 @@
1
  version https://git-lfs.github.com/spec/v1
2
- oid sha256:5d303985c595798bfeefb42ddd2e7161ff165594d8ef22c294bb33a042eff6ea
3
  size 3439
 
1
  version https://git-lfs.github.com/spec/v1
2
+ oid sha256:8269e66f591bb2ce05c5408c6e87e5e1915ae94de3850787ecd6fcf134a4ecce
3
  size 3439