obudzecie commited on Feb 27, 2024

Commit

10406c4

verified ·

1 Parent(s): 9d163bf

Training in progress, epoch 1

Browse files

Files changed (36) hide show

model.safetensors +1 -1
run-2/checkpoint-108/config.json +25 -0
run-2/checkpoint-108/model.safetensors +3 -0
run-2/checkpoint-108/optimizer.pt +3 -0
run-2/checkpoint-108/rng_state.pth +3 -0
run-2/checkpoint-108/scheduler.pt +3 -0
run-2/checkpoint-108/special_tokens_map.json +7 -0
run-2/checkpoint-108/tokenizer.json +0 -0
run-2/checkpoint-108/tokenizer_config.json +55 -0
run-2/checkpoint-108/trainer_state.json +44 -0
run-2/checkpoint-108/training_args.bin +3 -0
run-2/checkpoint-108/vocab.txt +0 -0
run-2/checkpoint-162/config.json +25 -0
run-2/checkpoint-162/model.safetensors +3 -0
run-2/checkpoint-162/optimizer.pt +3 -0
run-2/checkpoint-162/rng_state.pth +3 -0
run-2/checkpoint-162/scheduler.pt +3 -0
run-2/checkpoint-162/special_tokens_map.json +7 -0
run-2/checkpoint-162/tokenizer.json +0 -0
run-2/checkpoint-162/tokenizer_config.json +55 -0
run-2/checkpoint-162/trainer_state.json +53 -0
run-2/checkpoint-162/training_args.bin +3 -0
run-2/checkpoint-162/vocab.txt +0 -0
run-3/checkpoint-54/config.json +25 -0
run-3/checkpoint-54/model.safetensors +3 -0
run-3/checkpoint-54/optimizer.pt +3 -0
run-3/checkpoint-54/rng_state.pth +3 -0
run-3/checkpoint-54/scheduler.pt +3 -0
run-3/checkpoint-54/special_tokens_map.json +7 -0
run-3/checkpoint-54/tokenizer.json +0 -0
run-3/checkpoint-54/tokenizer_config.json +55 -0
run-3/checkpoint-54/trainer_state.json +35 -0
run-3/checkpoint-54/training_args.bin +3 -0
run-3/checkpoint-54/vocab.txt +0 -0
runs/Feb27_12-55-59_86e2730489af/events.out.tfevents.1709040044.86e2730489af.508.5 +3 -0
training_args.bin +1 -1

model.safetensors CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:14425fcaecc2d46a7e64c553c7c0e03a8b0b899a83ea0d3963fb80be859c7eb1
 size 267832560

 version https://git-lfs.github.com/spec/v1
+oid sha256:a07b3521cc8479b97635017fd08648a7bea2791ad6ec7bace28d41fb6d5321c4
 size 267832560

run-2/checkpoint-108/config.json ADDED Viewed

	@@ -0,0 +1,25 @@

+{
+  "_name_or_path": "distilbert-base-uncased",
+  "activation": "gelu",
+  "architectures": [
+    "DistilBertForSequenceClassification"
+  ],
+  "attention_dropout": 0.1,
+  "dim": 768,
+  "dropout": 0.1,
+  "hidden_dim": 3072,
+  "initializer_range": 0.02,
+  "max_position_embeddings": 512,
+  "model_type": "distilbert",
+  "n_heads": 12,
+  "n_layers": 6,
+  "pad_token_id": 0,
+  "problem_type": "single_label_classification",
+  "qa_dropout": 0.1,
+  "seq_classif_dropout": 0.2,
+  "sinusoidal_pos_embds": false,
+  "tie_weights_": true,
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.1",
+  "vocab_size": 30522
+}

run-2/checkpoint-108/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:94b6f1e3c1a543d3f1944a51003949b115ae8684a480dbb510f856e0cc2aa8fa
+size 267832560

run-2/checkpoint-108/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:bf9d74ca447919cb00c2eb2466bdbfc459d010f4ef881dabb2b5815b29de6ead
+size 535727290

run-2/checkpoint-108/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:455d3c21c1a460a690a7144ff5008874689ceac0c59ef21436ace505b2ab9be9
+size 14244

run-2/checkpoint-108/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:e2a64c6dc2736b6f206a7df5acdc80d91a138f9fc193e0cf29be85a539520c91
+size 1064

run-2/checkpoint-108/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "cls_token": "[CLS]",
+  "mask_token": "[MASK]",
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "unk_token": "[UNK]"
+}

run-2/checkpoint-108/tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

run-2/checkpoint-108/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,55 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "[PAD]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "100": {
+      "content": "[UNK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "101": {
+      "content": "[CLS]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "102": {
+      "content": "[SEP]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "103": {
+      "content": "[MASK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "clean_up_tokenization_spaces": true,
+  "cls_token": "[CLS]",
+  "do_lower_case": true,
+  "mask_token": "[MASK]",
+  "model_max_length": 512,
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "DistilBertTokenizer",
+  "unk_token": "[UNK]"
+}

run-2/checkpoint-108/trainer_state.json ADDED Viewed

	@@ -0,0 +1,44 @@

+{
+  "best_metric": 0.170590548044492,
+  "best_model_checkpoint": "distilbert-base-uncased-finetuned-cola/run-2/checkpoint-108",
+  "epoch": 2.0,
+  "eval_steps": 500,
+  "global_step": 108,
+  "is_hyper_param_search": true,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 1.0,
+      "eval_loss": 0.5949434041976929,
+      "eval_matthews_correlation": 0.0,
+      "eval_runtime": 0.6642,
+      "eval_samples_per_second": 1570.266,
+      "eval_steps_per_second": 99.365,
+      "step": 54
+    },
+    {
+      "epoch": 2.0,
+      "eval_loss": 0.5589540004730225,
+      "eval_matthews_correlation": 0.170590548044492,
+      "eval_runtime": 1.2474,
+      "eval_samples_per_second": 836.172,
+      "eval_steps_per_second": 52.912,
+      "step": 108
+    }
+  ],
+  "logging_steps": 500,
+  "max_steps": 162,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 3,
+  "save_steps": 500,
+  "total_flos": 0,
+  "train_batch_size": 16,
+  "trial_name": null,
+  "trial_params": {
+    "learning_rate": 1.8548433860537384e-05,
+    "num_train_epochs": 3,
+    "per_device_train_batch_size": 16,
+    "seed": 19
+  }
+}

run-2/checkpoint-108/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:f43db310149dc85ff8b52f32f702ad0e0fc7c0abc14ab704372d592275b95932
+size 4984

run-2/checkpoint-108/vocab.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

run-2/checkpoint-162/config.json ADDED Viewed

	@@ -0,0 +1,25 @@

+{
+  "_name_or_path": "distilbert-base-uncased",
+  "activation": "gelu",
+  "architectures": [
+    "DistilBertForSequenceClassification"
+  ],
+  "attention_dropout": 0.1,
+  "dim": 768,
+  "dropout": 0.1,
+  "hidden_dim": 3072,
+  "initializer_range": 0.02,
+  "max_position_embeddings": 512,
+  "model_type": "distilbert",
+  "n_heads": 12,
+  "n_layers": 6,
+  "pad_token_id": 0,
+  "problem_type": "single_label_classification",
+  "qa_dropout": 0.1,
+  "seq_classif_dropout": 0.2,
+  "sinusoidal_pos_embds": false,
+  "tie_weights_": true,
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.1",
+  "vocab_size": 30522
+}

run-2/checkpoint-162/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:c9467a19d971687c988e0b3a50436b0b97cfc04737df37bcf6c23b5fd2146d29
+size 267832560

run-2/checkpoint-162/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ab8022bf237f7049f70464034d79135f248abcf9117234843d9d16b30905adb5
+size 535727290

run-2/checkpoint-162/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:8d8342560780027bd62a65272e967a2d0fd106454080a61d45275e25df071c5e
+size 14244

run-2/checkpoint-162/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:65ee26e40b06c7695752cd5de136c5bdc13b2a7c6cd5f00a1063a9032f591bad
+size 1064

run-2/checkpoint-162/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "cls_token": "[CLS]",
+  "mask_token": "[MASK]",
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "unk_token": "[UNK]"
+}

run-2/checkpoint-162/tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

run-2/checkpoint-162/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,55 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "[PAD]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "100": {
+      "content": "[UNK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "101": {
+      "content": "[CLS]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "102": {
+      "content": "[SEP]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "103": {
+      "content": "[MASK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "clean_up_tokenization_spaces": true,
+  "cls_token": "[CLS]",
+  "do_lower_case": true,
+  "mask_token": "[MASK]",
+  "model_max_length": 512,
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "DistilBertTokenizer",
+  "unk_token": "[UNK]"
+}

run-2/checkpoint-162/trainer_state.json ADDED Viewed

	@@ -0,0 +1,53 @@

+{
+  "best_metric": 0.22258814800874943,
+  "best_model_checkpoint": "distilbert-base-uncased-finetuned-cola/run-2/checkpoint-162",
+  "epoch": 3.0,
+  "eval_steps": 500,
+  "global_step": 162,
+  "is_hyper_param_search": true,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 1.0,
+      "eval_loss": 0.5949434041976929,
+      "eval_matthews_correlation": 0.0,
+      "eval_runtime": 0.6642,
+      "eval_samples_per_second": 1570.266,
+      "eval_steps_per_second": 99.365,
+      "step": 54
+    },
+    {
+      "epoch": 2.0,
+      "eval_loss": 0.5589540004730225,
+      "eval_matthews_correlation": 0.170590548044492,
+      "eval_runtime": 1.2474,
+      "eval_samples_per_second": 836.172,
+      "eval_steps_per_second": 52.912,
+      "step": 108
+    },
+    {
+      "epoch": 3.0,
+      "eval_loss": 0.5674259066581726,
+      "eval_matthews_correlation": 0.22258814800874943,
+      "eval_runtime": 0.7895,
+      "eval_samples_per_second": 1321.051,
+      "eval_steps_per_second": 83.595,
+      "step": 162
+    }
+  ],
+  "logging_steps": 500,
+  "max_steps": 162,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 3,
+  "save_steps": 500,
+  "total_flos": 0,
+  "train_batch_size": 16,
+  "trial_name": null,
+  "trial_params": {
+    "learning_rate": 1.8548433860537384e-05,
+    "num_train_epochs": 3,
+    "per_device_train_batch_size": 16,
+    "seed": 19
+  }
+}

run-2/checkpoint-162/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:f43db310149dc85ff8b52f32f702ad0e0fc7c0abc14ab704372d592275b95932
+size 4984

run-2/checkpoint-162/vocab.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

run-3/checkpoint-54/config.json ADDED Viewed

	@@ -0,0 +1,25 @@

+{
+  "_name_or_path": "distilbert-base-uncased",
+  "activation": "gelu",
+  "architectures": [
+    "DistilBertForSequenceClassification"
+  ],
+  "attention_dropout": 0.1,
+  "dim": 768,
+  "dropout": 0.1,
+  "hidden_dim": 3072,
+  "initializer_range": 0.02,
+  "max_position_embeddings": 512,
+  "model_type": "distilbert",
+  "n_heads": 12,
+  "n_layers": 6,
+  "pad_token_id": 0,
+  "problem_type": "single_label_classification",
+  "qa_dropout": 0.1,
+  "seq_classif_dropout": 0.2,
+  "sinusoidal_pos_embds": false,
+  "tie_weights_": true,
+  "torch_dtype": "float32",
+  "transformers_version": "4.38.1",
+  "vocab_size": 30522
+}

run-3/checkpoint-54/model.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:a07b3521cc8479b97635017fd08648a7bea2791ad6ec7bace28d41fb6d5321c4
+size 267832560

run-3/checkpoint-54/optimizer.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d3c9bdaa0c6dd042796f0553c74172060ba3f9ff26517f16579571c24bcf23ae
+size 535727290

run-3/checkpoint-54/rng_state.pth ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:5b760589f8554fc9a823461802d959d20ee6933034d468005d9e9650985646e1
+size 14244

run-3/checkpoint-54/scheduler.pt ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6e821cd8c0cea6045dab45710ac190a0970d06aa7c327d90cc1165ca399b6078
+size 1064

run-3/checkpoint-54/special_tokens_map.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "cls_token": "[CLS]",
+  "mask_token": "[MASK]",
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "unk_token": "[UNK]"
+}

run-3/checkpoint-54/tokenizer.json ADDED Viewed

The diff for this file is too large to render. See raw diff

run-3/checkpoint-54/tokenizer_config.json ADDED Viewed

	@@ -0,0 +1,55 @@

+{
+  "added_tokens_decoder": {
+    "0": {
+      "content": "[PAD]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "100": {
+      "content": "[UNK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "101": {
+      "content": "[CLS]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "102": {
+      "content": "[SEP]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "103": {
+      "content": "[MASK]",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "clean_up_tokenization_spaces": true,
+  "cls_token": "[CLS]",
+  "do_lower_case": true,
+  "mask_token": "[MASK]",
+  "model_max_length": 512,
+  "pad_token": "[PAD]",
+  "sep_token": "[SEP]",
+  "strip_accents": null,
+  "tokenize_chinese_chars": true,
+  "tokenizer_class": "DistilBertTokenizer",
+  "unk_token": "[UNK]"
+}

run-3/checkpoint-54/trainer_state.json ADDED Viewed

	@@ -0,0 +1,35 @@

+{
+  "best_metric": 0.0,
+  "best_model_checkpoint": "distilbert-base-uncased-finetuned-cola/run-3/checkpoint-54",
+  "epoch": 1.0,
+  "eval_steps": 500,
+  "global_step": 54,
+  "is_hyper_param_search": true,
+  "is_local_process_zero": true,
+  "is_world_process_zero": true,
+  "log_history": [
+    {
+      "epoch": 1.0,
+      "eval_loss": 0.6135343909263611,
+      "eval_matthews_correlation": 0.0,
+      "eval_runtime": 0.6686,
+      "eval_samples_per_second": 1560.01,
+      "eval_steps_per_second": 98.716,
+      "step": 54
+    }
+  ],
+  "logging_steps": 500,
+  "max_steps": 108,
+  "num_input_tokens_seen": 0,
+  "num_train_epochs": 2,
+  "save_steps": 500,
+  "total_flos": 0,
+  "train_batch_size": 16,
+  "trial_name": null,
+  "trial_params": {
+    "learning_rate": 7.371411848219159e-06,
+    "num_train_epochs": 2,
+    "per_device_train_batch_size": 16,
+    "seed": 35
+  }
+}

run-3/checkpoint-54/training_args.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:c6ae447073f49fbe1616e5d8dd0fa3e8b79bcfd527890bf6b4740cb1b9964bfc
+size 4984

run-3/checkpoint-54/vocab.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

runs/Feb27_12-55-59_86e2730489af/events.out.tfevents.1709040044.86e2730489af.508.5 ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d2ea12e0669b8505d131ae8858651dd1a64c97adab2e934265c62ff90b93dd07
+size 5548

training_args.bin CHANGED Viewed

@@ -1,3 +1,3 @@
 version https://git-lfs.github.com/spec/v1
-oid sha256:f43db310149dc85ff8b52f32f702ad0e0fc7c0abc14ab704372d592275b95932
 size 4984

 version https://git-lfs.github.com/spec/v1
+oid sha256:c6ae447073f49fbe1616e5d8dd0fa3e8b79bcfd527890bf6b4740cb1b9964bfc
 size 4984