sst2 / training_config.json
ShengdingHu's picture
Training in progress, step 200
b81d89e
{"dataset_config_name": ["en"], "delta_type": "lora", "do_eval": true, "do_test": true, "do_train": true, "eval_dataset_config_name": ["en"], "eval_dataset_name": "sst2", "eval_steps": 200, "evaluation_strategy": "steps", "greater_is_better": true, "is_seq2seq": false, "learning_rate": 0.001, "load_best_model_at_end": true, "max_source_length": 128, "metric_for_best_model": "average_metrics", "model_name_or_path": "/home/hushengding/plm_cache/bigbird-roberta-large", "modified_modules": ["query", "key"], "num_train_epochs": 3, "output_dir": "outputs/lora/bigbird-roberta-large/sst2", "overwrite_output_dir": true, "per_device_eval_batch_size": 32, "per_device_train_batch_size": 32, "predict_with_generate": false, "push_to_hub": true, "save_steps": 200, "save_strategy": "steps", "save_total_limit": 1, "seed": 42, "split_validation_test": true, "task_name": "sst2", "test_dataset_config_name": ["en"], "test_dataset_name": "sst2", "tokenizer_name": "/home/hushengding/plm_cache/bigbird-roberta-large", "warmup_steps": 0}