aisingapore
/

llama3-8b-cpt-sea-lionv2.1-instruct

Text Generation

text-generation-inference

Inference Endpoints

Model card Files Files and versions Community

xianbin commited on Jul 30, 2024

Commit

80baaa0

·

verified ·

1 Parent(s): d60401a

Update config.json

Files changed (1) hide show

config.json +1 -3

config.json CHANGED Viewed

@@ -12,7 +12,6 @@
   "intermediate_size": 14336,
   "max_position_embeddings": 8192,
   "model_type": "llama",
-  "name": "hf_causal_lm",
   "num_attention_heads": 32,
   "num_hidden_layers": 32,
   "num_key_value_heads": 8,
@@ -21,9 +20,8 @@
   "rope_scaling": null,
   "rope_theta": 500000.0,
   "tie_word_embeddings": false,
-  "torch_dtype": "float32",
   "transformers_version": "4.43.2",
-  "trust_remote_code": true,
   "use_cache": true,
   "use_flash_attention_2": true,
   "vocab_size": 128256

   "intermediate_size": 14336,
   "max_position_embeddings": 8192,
   "model_type": "llama",
   "num_attention_heads": 32,
   "num_hidden_layers": 32,
   "num_key_value_heads": 8,
   "rope_scaling": null,
   "rope_theta": 500000.0,
   "tie_word_embeddings": false,
+  "torch_dtype": "bfloat16",
   "transformers_version": "4.43.2",
   "use_cache": true,
   "use_flash_attention_2": true,
   "vocab_size": 128256