![finnstrom3693's picture](https://cdn-avatars.huggingface.co/v1/production/uploads/noauth/R15ksyXv72eZttRtsgMvA.png)
training using pytorch native 3 epoch, batch size 14, block size 512,lr 1e-4 cosine
8d45360
verified
{ | |
"add_prefix_space": false, | |
"added_tokens_decoder": { | |
"50256": { | |
"content": "<|endoftext|>", | |
"lstrip": false, | |
"normalized": true, | |
"rstrip": false, | |
"single_word": false, | |
"special": true | |
} | |
}, | |
"bos_token": "<|endoftext|>", | |
"clean_up_tokenization_spaces": true, | |
"eos_token": "<|endoftext|>", | |
"model_max_length": 1024, | |
"pad_token": "<|endoftext|>", | |
"tokenizer_class": "GPT2Tokenizer", | |
"unk_token": "<|endoftext|>" | |
} | |