Training in progress, step 500

Files changed (11) hide show

.gitignore ADDED Viewed

	@@ -0,0 +1 @@


1	+ checkpoint-*/

added_tokens.json ADDED Viewed

+{
+  "[CLAIM_END]": 128006,
+  "[CLAIM_START]": 128005,
+  "[CONCLUDING STATEMENT_END]": 128014,
+  "[CONCLUDING STATEMENT_START]": 128013,
+  "[COUNTERCLAIM_END]": 128008,
+  "[COUNTERCLAIM_START]": 128007,
+  "[EVIDENCE_END]": 128012,
+  "[EVIDENCE_START]": 128011,
+  "[LEAD_END]": 128002,
+  "[LEAD_START]": 128001,
+  "[MASK]": 128000,
+  "[POSITION_END]": 128004,
+  "[POSITION_START]": 128003,
+  "[REBUTTAL_END]": 128010,
+  "[REBUTTAL_START]": 128009
+}

config.json ADDED Viewed

+{
+  "_name_or_path": "microsoft/deberta-v3-large",
+  "architectures": [
+    "DebertaV2ForMaskedLM"
+  ],
+  "attention_probs_dropout_prob": 0.1,
+  "hidden_act": "gelu",
+  "hidden_dropout_prob": 0.1,
+  "hidden_size": 1024,
+  "initializer_range": 0.02,
+  "intermediate_size": 4096,
+  "layer_norm_eps": 1e-07,
+  "max_position_embeddings": 1024,
+  "max_relative_positions": 128,
+  "model_type": "deberta-v2",
+  "norm_rel_ebd": "layer_norm",
+  "num_attention_heads": 16,
+  "num_hidden_layers": 24,
+  "pad_token_id": 0,
+  "pooler_dropout": 0,
+  "pooler_hidden_act": "gelu",
+  "pooler_hidden_size": 1024,
+  "pos_att_type": [
+    "p2c",
+    "c2p"
+  ],
+  "position_biased_input": false,
+  "position_buckets": 256,
+  "relative_attention": true,
+  "share_att_key": true,
+  "torch_dtype": "float32",
+  "transformers_version": "4.24.0",
+  "type_vocab_size": 0,
+  "vocab_size": 128100
+}

pytorch_model.bin ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:3c89ad41093f771147c93ac9409920434a2afe4d1bf4ad54ef8a8eab973d126b
+size 1740910441

runs/Nov17_19-41-07_21a669a84046/1668714087.2024877/events.out.tfevents.1668714087.21a669a84046.16332.1 ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:3dffc5aef45258a6eefd3925975786ea78f8d29e9a6e29cf0b3ebe1a6561d940
+size 5493

runs/Nov17_19-41-07_21a669a84046/events.out.tfevents.1668714087.21a669a84046.16332.0 ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:7060b217e39a5e5ddc765346b9eeaba9fe0ba5a0930f901b3ab0b058a9863042
+size 4352

special_tokens_map.json ADDED Viewed

+{
+  "additional_special_tokens": [
+    "[LEAD_START]",
+    "[LEAD_END]",
+    "[POSITION_START]",
+    "[POSITION_END]",
+    "[CLAIM_START]",
+    "[CLAIM_END]",
+    "[COUNTERCLAIM_START]",
+    "[COUNTERCLAIM_END]",
+    "[REBUTTAL_START]",
+    "[REBUTTAL_END]",
+    "[EVIDENCE_START]",
+    "[EVIDENCE_END]",
+    "[CONCLUDING STATEMENT_START]",
+    "[CONCLUDING STATEMENT_END]"
+  ],
+  "bos_token": "[CLS]",
+  "cls_token": "[CLS]",
+  "eos_token": "[SEP]",
+  "mask_token": "[MASK]",
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "unk_token": "[UNK]"
+}

spm.model ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:c679fbf93643d19aab7ee10c0b99e460bdbc02fedf34b92b05af343b4af586fd
+size 2464616

tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

tokenizer_config.json ADDED Viewed

+{
+  "additional_special_tokens": [
+    "[LEAD_START]",
+    "[LEAD_END]",
+    "[POSITION_START]",
+    "[POSITION_END]",
+    "[CLAIM_START]",
+    "[CLAIM_END]",
+    "[COUNTERCLAIM_START]",
+    "[COUNTERCLAIM_END]",
+    "[REBUTTAL_START]",
+    "[REBUTTAL_END]",
+    "[EVIDENCE_START]",
+    "[EVIDENCE_END]",
+    "[CONCLUDING STATEMENT_START]",
+    "[CONCLUDING STATEMENT_END]"
+  ],
+  "bos_token": "[CLS]",
+  "cls_token": "[CLS]",
+  "do_lower_case": false,
+  "eos_token": "[SEP]",
+  "mask_token": "[MASK]",
+  "name_or_path": "microsoft/deberta-v3-large",
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "sp_model_kwargs": {},
+  "special_tokens_map_file": null,
+  "split_by_punct": false,
+  "tokenizer_class": "DebertaV2Tokenizer",
+  "unk_token": "[UNK]",
+  "vocab_type": "spm"
+}

training_args.bin ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:444a562a2932e11531857c5f5f67c1f9cd21d7a7b2be15e0eb3b8da0a94f7c56
+size 3375