Upload folder using huggingface_hub

Files changed (4) hide show

README.md CHANGED Viewed

@@ -1,6 +1,6 @@
 ---
-library_name: transformers
-tags: []
 ---
 # Model Card for Model ID
@@ -15,7 +15,7 @@ tags: []
 <!-- Provide a longer summary of what this model is. -->
-This is the model card of a 🤗 transformers model that has been pushed on the Hub. This model card has been automatically generated.
 - **Developed by:** [More Information Needed]
 - **Funded by [optional]:** [More Information Needed]
@@ -196,4 +196,7 @@ Carbon emissions can be estimated using the [Machine Learning Impact calculator]
 ## Model Card Contact
-[More Information Needed]

 ---
+base_model: CohereForAI/aya-23-8B
+library_name: peft
 ---
 # Model Card for Model ID
 <!-- Provide a longer summary of what this model is. -->
 - **Developed by:** [More Information Needed]
 - **Funded by [optional]:** [More Information Needed]
 ## Model Card Contact
+[More Information Needed]
+### Framework versions
+- PEFT 0.12.0

config.json ADDED Viewed

+{
+  "_name_or_path": "CohereForAI/aya-23-8B",
+  "architectures": [
+    "CohereForCausalLM"
+  ],
+  "attention_bias": false,
+  "attention_dropout": 0.0,
+  "bos_token_id": 5,
+  "eos_token_id": 255001,
+  "hidden_act": "silu",
+  "hidden_size": 4096,
+  "initializer_range": 0.02,
+  "intermediate_size": 14336,
+  "layer_norm_eps": 1e-05,
+  "logit_scale": 0.0625,
+  "max_position_embeddings": 8192,
+  "model_type": "cohere",
+  "num_attention_heads": 32,
+  "num_hidden_layers": 32,
+  "num_key_value_heads": 8,
+  "pad_token_id": 0,
+  "rope_theta": 10000,
+  "torch_dtype": "float16",
+  "transformers_version": "4.42.4",
+  "use_cache": true,
+  "use_qk_norm": false,
+  "vocab_size": 256000
+}

pytorch_model.bin ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:681d6dcf0f43ca3078bc6b9b8696ef83359d37175ae4f1c5e1360f6f17ea6fea
+size 32125992678

training_args.bin ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:f12cfd1e48b83f8c4886e20ae55741d97af93ee5df31310622a9c258fda16b78
+size 5048