pkairpsp commited on 11 days ago

Commit

b7ce9b8

verified ·

1 Parent(s): 82eb079

Upload folder using huggingface_hub

Browse files

Files changed (33) hide show

checkpoint-238/config.json +36 -0
checkpoint-238/model.safetensors +3 -0
checkpoint-238/optimizer.pt +3 -0
checkpoint-238/rng_state.pth +3 -0
checkpoint-238/scheduler.pt +3 -0
checkpoint-238/special_tokens_map.json +37 -0
checkpoint-238/tokenizer.json +0 -0
checkpoint-238/tokenizer_config.json +65 -0
checkpoint-238/trainer_state.json +45 -0
checkpoint-238/training_args.bin +3 -0
checkpoint-238/vocab.txt +0 -0
checkpoint-476/config.json +36 -0
checkpoint-476/model.safetensors +3 -0
checkpoint-476/optimizer.pt +3 -0
checkpoint-476/rng_state.pth +3 -0
checkpoint-476/scheduler.pt +3 -0
checkpoint-476/special_tokens_map.json +37 -0
checkpoint-476/tokenizer.json +0 -0
checkpoint-476/tokenizer_config.json +65 -0
checkpoint-476/trainer_state.json +57 -0
checkpoint-476/training_args.bin +3 -0
checkpoint-476/vocab.txt +0 -0
checkpoint-714/config.json +36 -0
checkpoint-714/model.safetensors +3 -0
checkpoint-714/optimizer.pt +3 -0
checkpoint-714/rng_state.pth +3 -0
checkpoint-714/scheduler.pt +3 -0
checkpoint-714/special_tokens_map.json +37 -0
checkpoint-714/tokenizer.json +0 -0
checkpoint-714/tokenizer_config.json +65 -0
checkpoint-714/trainer_state.json +76 -0
checkpoint-714/training_args.bin +3 -0
checkpoint-714/vocab.txt +0 -0

checkpoint-238/config.json ADDED Viewed

	@@ -0,0 +1,36 @@

+{
+  "_name_or_path": "lxyuan/distilbert-base-multilingual-cased-sentiments-student",
+  "activation": "gelu",
+  "architectures": [
+    "DistilBertForSequenceClassification"
+  ],
+  "attention_dropout": 0.1,
+  "dim": 768,
+  "dropout": 0.1,
+  "hidden_dim": 3072,
+  "id2label": {
+    "0": "positive",
+    "1": "neutral",
+    "2": "negative"
+  },
+  "initializer_range": 0.02,
+  "label2id": {
+    "negative": 2,
+    "neutral": 1,
+    "positive": 0
+  },
+  "max_position_embeddings": 512,
+  "model_type": "distilbert",
+  "n_heads": 12,
+  "n_layers": 6,
+  "output_past": true,
+  "pad_token_id": 0,
+  "problem_type": "single_label_classification",
+  "qa_dropout": 0.1,
+  "seq_classif_dropout": 0.2,
+  "sinusoidal_pos_embds": false,
+  "tie_weights_": true,
+  "torch_dtype": "float32",
+  "transformers_version": "4.48.2",
+  "vocab_size": 119547
+}

checkpoint-238/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:1799e7ac41b455ae6f71560b82c6855d2d2e3c456c104b73e8eddc3407106245
+size 541320452

checkpoint-238/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:616b70eeecaef9ab77a91276c1fc19a4cc03cae2bbd8bde7fb31d70b40fd43a0
+size 1082700154

checkpoint-238/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6b452c3116d0684fa8abb970ebeb744784ffe3b664f3b2a67d4d1427fa15c23f
+size 13990

checkpoint-238/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:41bdca404a52091ce9b3b05d7e5b9e8fbc85c2940dd44523ae94aa1058f7cfdc
+size 1064

checkpoint-238/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,37 @@

+{
+  "cls_token": {
+    "content": "[CLS]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "mask_token": {
+    "content": "[MASK]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": {
+    "content": "[PAD]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "sep_token": {
+    "content": "[SEP]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "unk_token": {
+    "content": "[UNK]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

checkpoint-238/tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

checkpoint-238/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,65 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "[PAD]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "100": {
+      "content": "[UNK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "101": {
+      "content": "[CLS]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "102": {
+      "content": "[SEP]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "103": {
+      "content": "[MASK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "clean_up_tokenization_spaces": true,
+  "cls_token": "[CLS]",
+  "do_basic_tokenize": true,
+  "do_lower_case": false,
+  "extra_special_tokens": {},
+  "mask_token": "[MASK]",
+  "max_length": 512,
+  "model_max_length": 512,
+  "never_split": null,
+  "pad_to_multiple_of": null,
+  "pad_token": "[PAD]",
+  "pad_token_type_id": 0,
+  "padding_side": "right",
+  "sep_token": "[SEP]",
+  "stride": 0,
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "DistilBertTokenizer",
+  "truncation_side": "right",
+  "truncation_strategy": "longest_first",
+  "unk_token": "[UNK]"
+}

checkpoint-238/trainer_state.json ADDED Viewed

	@@ -0,0 +1,45 @@

+{
+  "best_metric": null,
+  "best_model_checkpoint": null,
+  "epoch": 1.0,
+  "eval_steps": 500,
+  "global_step": 238,
+  "is_hyper_param_search": false,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 1.0,
+      "eval_accuracy": 0.9245283018867925,
+      "eval_f1": 0.9242719879430212,
+      "eval_loss": 0.25112205743789673,
+      "eval_precision": 0.9315685227148311,
+      "eval_recall": 0.9245283018867925,
+      "eval_runtime": 292.267,
+      "eval_samples_per_second": 1.088,
+      "eval_steps_per_second": 0.137,
+      "step": 238
+    }
+  ],
+  "logging_steps": 500,
+  "max_steps": 714,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 3,
+  "save_steps": 500,
+  "stateful_callbacks": {
+    "TrainerControl": {
+      "args": {
+        "should_epoch_stop": false,
+        "should_evaluate": false,
+        "should_log": false,
+        "should_save": true,
+        "should_training_stop": false
+      },
+      "attributes": {}
+    }
+  },
+  "total_flos": 251560076184576.0,
+  "train_batch_size": 8,
+  "trial_name": null,
+  "trial_params": null
+}

checkpoint-238/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:86abd7c46c1196bbbf8733498652cac4d01b411776297fbb620a2b871cbfda9b
+size 5304

checkpoint-238/vocab.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

checkpoint-476/config.json ADDED Viewed

	@@ -0,0 +1,36 @@

+{
+  "_name_or_path": "lxyuan/distilbert-base-multilingual-cased-sentiments-student",
+  "activation": "gelu",
+  "architectures": [
+    "DistilBertForSequenceClassification"
+  ],
+  "attention_dropout": 0.1,
+  "dim": 768,
+  "dropout": 0.1,
+  "hidden_dim": 3072,
+  "id2label": {
+    "0": "positive",
+    "1": "neutral",
+    "2": "negative"
+  },
+  "initializer_range": 0.02,
+  "label2id": {
+    "negative": 2,
+    "neutral": 1,
+    "positive": 0
+  },
+  "max_position_embeddings": 512,
+  "model_type": "distilbert",
+  "n_heads": 12,
+  "n_layers": 6,
+  "output_past": true,
+  "pad_token_id": 0,
+  "problem_type": "single_label_classification",
+  "qa_dropout": 0.1,
+  "seq_classif_dropout": 0.2,
+  "sinusoidal_pos_embds": false,
+  "tie_weights_": true,
+  "torch_dtype": "float32",
+  "transformers_version": "4.48.2",
+  "vocab_size": 119547
+}

checkpoint-476/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:3c33c5215bed4e231aa2ea521ffef4cb8cbd23996af6b45b1bcff2b914dc47c9
+size 541320452

checkpoint-476/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6301ea64791855a023c8591e4ca24279dcf2ed3a6427b5497270d0c3d32bcb62
+size 1082700154

checkpoint-476/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:21e9377de387e81da5641d438c58da652edc70e4e1054413c019370365b92668
+size 13990

checkpoint-476/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:58afa012b4af9816e109a7c268738f37084e6bf720c5b16e7170a3740582f1f2
+size 1064

checkpoint-476/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,37 @@

+{
+  "cls_token": {
+    "content": "[CLS]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "mask_token": {
+    "content": "[MASK]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": {
+    "content": "[PAD]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "sep_token": {
+    "content": "[SEP]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "unk_token": {
+    "content": "[UNK]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

checkpoint-476/tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

checkpoint-476/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,65 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "[PAD]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "100": {
+      "content": "[UNK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "101": {
+      "content": "[CLS]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "102": {
+      "content": "[SEP]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "103": {
+      "content": "[MASK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "clean_up_tokenization_spaces": true,
+  "cls_token": "[CLS]",
+  "do_basic_tokenize": true,
+  "do_lower_case": false,
+  "extra_special_tokens": {},
+  "mask_token": "[MASK]",
+  "max_length": 512,
+  "model_max_length": 512,
+  "never_split": null,
+  "pad_to_multiple_of": null,
+  "pad_token": "[PAD]",
+  "pad_token_type_id": 0,
+  "padding_side": "right",
+  "sep_token": "[SEP]",
+  "stride": 0,
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "DistilBertTokenizer",
+  "truncation_side": "right",
+  "truncation_strategy": "longest_first",
+  "unk_token": "[UNK]"
+}

checkpoint-476/trainer_state.json ADDED Viewed

	@@ -0,0 +1,57 @@

+{
+  "best_metric": null,
+  "best_model_checkpoint": null,
+  "epoch": 2.0,
+  "eval_steps": 500,
+  "global_step": 476,
+  "is_hyper_param_search": false,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 1.0,
+      "eval_accuracy": 0.9245283018867925,
+      "eval_f1": 0.9242719879430212,
+      "eval_loss": 0.25112205743789673,
+      "eval_precision": 0.9315685227148311,
+      "eval_recall": 0.9245283018867925,
+      "eval_runtime": 292.267,
+      "eval_samples_per_second": 1.088,
+      "eval_steps_per_second": 0.137,
+      "step": 238
+    },
+    {
+      "epoch": 2.0,
+      "eval_accuracy": 0.9339622641509434,
+      "eval_f1": 0.933657504155167,
+      "eval_loss": 0.14069578051567078,
+      "eval_precision": 0.9421693229583453,
+      "eval_recall": 0.9339622641509434,
+      "eval_runtime": 289.7835,
+      "eval_samples_per_second": 1.097,
+      "eval_steps_per_second": 0.138,
+      "step": 476
+    }
+  ],
+  "logging_steps": 500,
+  "max_steps": 714,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 3,
+  "save_steps": 500,
+  "stateful_callbacks": {
+    "TrainerControl": {
+      "args": {
+        "should_epoch_stop": false,
+        "should_evaluate": false,
+        "should_log": false,
+        "should_save": true,
+        "should_training_stop": false
+      },
+      "attributes": {}
+    }
+  },
+  "total_flos": 503120152369152.0,
+  "train_batch_size": 8,
+  "trial_name": null,
+  "trial_params": null
+}

checkpoint-476/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:86abd7c46c1196bbbf8733498652cac4d01b411776297fbb620a2b871cbfda9b
+size 5304

checkpoint-476/vocab.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

checkpoint-714/config.json ADDED Viewed

	@@ -0,0 +1,36 @@

+{
+  "_name_or_path": "lxyuan/distilbert-base-multilingual-cased-sentiments-student",
+  "activation": "gelu",
+  "architectures": [
+    "DistilBertForSequenceClassification"
+  ],
+  "attention_dropout": 0.1,
+  "dim": 768,
+  "dropout": 0.1,
+  "hidden_dim": 3072,
+  "id2label": {
+    "0": "positive",
+    "1": "neutral",
+    "2": "negative"
+  },
+  "initializer_range": 0.02,
+  "label2id": {
+    "negative": 2,
+    "neutral": 1,
+    "positive": 0
+  },
+  "max_position_embeddings": 512,
+  "model_type": "distilbert",
+  "n_heads": 12,
+  "n_layers": 6,
+  "output_past": true,
+  "pad_token_id": 0,
+  "problem_type": "single_label_classification",
+  "qa_dropout": 0.1,
+  "seq_classif_dropout": 0.2,
+  "sinusoidal_pos_embds": false,
+  "tie_weights_": true,
+  "torch_dtype": "float32",
+  "transformers_version": "4.48.2",
+  "vocab_size": 119547
+}

checkpoint-714/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:7c0ff7996a4877cee32eb9bfacd2aaab4a9f5a6a70ca301d119bc6f20431bdb5
+size 541320452

checkpoint-714/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:1394f230e2e499c85b6362d65139a57c17d4fab9128d8f239034d86bd1d05d60
+size 1082700154

checkpoint-714/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:23550270f41dd9b1b88f47c15a0ab60a6fea0d16274d0b4e3d2908b4138d0456
+size 13990

checkpoint-714/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:15697dbb71dffb8de0632092db9d80e708ce9b665d10b2108a5674c6a0471210
+size 1064

checkpoint-714/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,37 @@

+{
+  "cls_token": {
+    "content": "[CLS]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "mask_token": {
+    "content": "[MASK]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": {
+    "content": "[PAD]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "sep_token": {
+    "content": "[SEP]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "unk_token": {
+    "content": "[UNK]",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

checkpoint-714/tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

checkpoint-714/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,65 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "[PAD]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "100": {
+      "content": "[UNK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "101": {
+      "content": "[CLS]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "102": {
+      "content": "[SEP]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "103": {
+      "content": "[MASK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "clean_up_tokenization_spaces": true,
+  "cls_token": "[CLS]",
+  "do_basic_tokenize": true,
+  "do_lower_case": false,
+  "extra_special_tokens": {},
+  "mask_token": "[MASK]",
+  "max_length": 512,
+  "model_max_length": 512,
+  "never_split": null,
+  "pad_to_multiple_of": null,
+  "pad_token": "[PAD]",
+  "pad_token_type_id": 0,
+  "padding_side": "right",
+  "sep_token": "[SEP]",
+  "stride": 0,
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "DistilBertTokenizer",
+  "truncation_side": "right",
+  "truncation_strategy": "longest_first",
+  "unk_token": "[UNK]"
+}

checkpoint-714/trainer_state.json ADDED Viewed

	@@ -0,0 +1,76 @@

+{
+  "best_metric": null,
+  "best_model_checkpoint": null,
+  "epoch": 3.0,
+  "eval_steps": 500,
+  "global_step": 714,
+  "is_hyper_param_search": false,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 1.0,
+      "eval_accuracy": 0.9245283018867925,
+      "eval_f1": 0.9242719879430212,
+      "eval_loss": 0.25112205743789673,
+      "eval_precision": 0.9315685227148311,
+      "eval_recall": 0.9245283018867925,
+      "eval_runtime": 292.267,
+      "eval_samples_per_second": 1.088,
+      "eval_steps_per_second": 0.137,
+      "step": 238
+    },
+    {
+      "epoch": 2.0,
+      "eval_accuracy": 0.9339622641509434,
+      "eval_f1": 0.933657504155167,
+      "eval_loss": 0.14069578051567078,
+      "eval_precision": 0.9421693229583453,
+      "eval_recall": 0.9339622641509434,
+      "eval_runtime": 289.7835,
+      "eval_samples_per_second": 1.097,
+      "eval_steps_per_second": 0.138,
+      "step": 476
+    },
+    {
+      "epoch": 2.100840336134454,
+      "grad_norm": 0.010930689983069897,
+      "learning_rate": 5.994397759103642e-06,
+      "loss": 0.1044,
+      "step": 500
+    },
+    {
+      "epoch": 3.0,
+      "eval_accuracy": 0.9308176100628931,
+      "eval_f1": 0.9305292147163376,
+      "eval_loss": 0.3272361159324646,
+      "eval_precision": 0.9385589054841049,
+      "eval_recall": 0.9308176100628931,
+      "eval_runtime": 294.9296,
+      "eval_samples_per_second": 1.078,
+      "eval_steps_per_second": 0.136,
+      "step": 714
+    }
+  ],
+  "logging_steps": 500,
+  "max_steps": 714,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 3,
+  "save_steps": 500,
+  "stateful_callbacks": {
+    "TrainerControl": {
+      "args": {
+        "should_epoch_stop": false,
+        "should_evaluate": false,
+        "should_log": false,
+        "should_save": true,
+        "should_training_stop": true
+      },
+      "attributes": {}
+    }
+  },
+  "total_flos": 754680228553728.0,
+  "train_batch_size": 8,
+  "trial_name": null,
+  "trial_params": null
+}

checkpoint-714/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:86abd7c46c1196bbbf8733498652cac4d01b411776297fbb620a2b871cbfda9b
+size 5304

checkpoint-714/vocab.txt ADDED Viewed

The diff for this file is too large to render. See raw diff