s3nh commited on
Commit
9e0808b
1 Parent(s): be35290

Upload ./ with huggingface_hub

Browse files
.gitattributes CHANGED
@@ -33,3 +33,13 @@ saved_model/**/* filter=lfs diff=lfs merge=lfs -text
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
 
 
 
 
 
 
 
 
 
 
 
33
  *.zip filter=lfs diff=lfs merge=lfs -text
34
  *.zst filter=lfs diff=lfs merge=lfs -text
35
  *tfevents* filter=lfs diff=lfs merge=lfs -text
36
+ Sentdex-WSB-GPT-13B.Q6_K.gguf filter=lfs diff=lfs merge=lfs -text
37
+ Sentdex-WSB-GPT-13B.Q3_K_S.gguf filter=lfs diff=lfs merge=lfs -text
38
+ Sentdex-WSB-GPT-13B.Q5_K_S.gguf filter=lfs diff=lfs merge=lfs -text
39
+ Sentdex-WSB-GPT-13B.Q3_K_M.gguf filter=lfs diff=lfs merge=lfs -text
40
+ Sentdex-WSB-GPT-13B.Q4_K_S.gguf filter=lfs diff=lfs merge=lfs -text
41
+ Sentdex-WSB-GPT-13B.Q2_K.gguf filter=lfs diff=lfs merge=lfs -text
42
+ Sentdex-WSB-GPT-13B.Q4_K_M.gguf filter=lfs diff=lfs merge=lfs -text
43
+ Sentdex-WSB-GPT-13B.Q3_K_L.gguf filter=lfs diff=lfs merge=lfs -text
44
+ Sentdex-WSB-GPT-13B.Q8_0.gguf filter=lfs diff=lfs merge=lfs -text
45
+ Sentdex-WSB-GPT-13B.Q5_K_M.gguf filter=lfs diff=lfs merge=lfs -text
README.md ADDED
@@ -0,0 +1,63 @@
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
 
1
+
2
+ ---
3
+ license: openrail
4
+ pipeline_tag: text-generation
5
+ library_name: transformers
6
+ language:
7
+ - zh
8
+ - en
9
+ ---
10
+
11
+
12
+ ## Original model card
13
+
14
+ Buy me a coffee if you like this project ;)
15
+ <a href="https://www.buymeacoffee.com/s3nh"><img src="https://www.buymeacoffee.com/assets/img/guidelines/download-assets-sm-1.svg" alt=""></a>
16
+
17
+ #### Description
18
+
19
+ GGUF Format model files for [This project](https://huggingface.co/Sentdex/WSB-GPT-13B).
20
+
21
+ ### GGUF Specs
22
+
23
+ GGUF is a format based on the existing GGJT, but makes a few changes to the format to make it more extensible and easier to use. The following features are desired:
24
+
25
+ Single-file deployment: they can be easily distributed and loaded, and do not require any external files for additional information.
26
+ Extensible: new features can be added to GGML-based executors/new information can be added to GGUF models without breaking compatibility with existing models.
27
+ mmap compatibility: models can be loaded using mmap for fast loading and saving.
28
+ Easy to use: models can be easily loaded and saved using a small amount of code, with no need for external libraries, regardless of the language used.
29
+ Full information: all information needed to load a model is contained in the model file, and no additional information needs to be provided by the user.
30
+ The key difference between GGJT and GGUF is the use of a key-value structure for the hyperparameters (now referred to as metadata), rather than a list of untyped values.
31
+ This allows for new metadata to be added without breaking compatibility with existing models, and to annotate the model with additional information that may be useful for
32
+ inference or for identifying the model.
33
+
34
+ ### Perplexity params
35
+
36
+ Model Measure Q2_K Q3_K_S Q3_K_M Q3_K_L Q4_0 Q4_1 Q4_K_S Q4_K_M Q5_0 Q5_1 Q5_K_S Q5_K_M Q6_K Q8_0 F16
37
+ 7B perplexity 6.7764 6.4571 6.1503 6.0869 6.1565 6.0912 6.0215 5.9601 5.9862 5.9481 5.9419 5.9208 5.9110 5.9070 5.9066
38
+ 13B perplexity 5.8545 5.6033 5.4498 5.4063 5.3860 5.3608 5.3404 5.3002 5.2856 5.2706 5.2785 5.2638 5.2568 5.2548 5.2543
39
+
40
+
41
+
42
+ ### inference
43
+
44
+
45
+ ```python
46
+
47
+ import ctransformers
48
+
49
+ from ctransformers import AutoModelForCausalLM
50
+
51
+ model = AutoModelForCausalLM.from_pretrained(output_dir, gguf_file,
52
+ gpu_layers=32, model_type="llama")
53
+
54
+ manual_input: str = "Tell me about your last dream, please."
55
+
56
+ llm(manual_input,
57
+ max_new_tokens=256,
58
+ temperature=0.9,
59
+ top_p= 0.7)
60
+
61
+ ```
62
+
63
+ # Original model card
Sentdex-WSB-GPT-13B.Q2_K.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:e8f40d1ffc4541886b531647500d9bbc97cf992fee0cc9a467a3484a7ca5dfc0
3
+ size 5429348192
Sentdex-WSB-GPT-13B.Q3_K_L.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:0639fb8ddad92297ae65f51340b9ef7b1bed65d882431d0f892198eee108e09a
3
+ size 6929559392
Sentdex-WSB-GPT-13B.Q3_K_M.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:9edc00f958be175a47e29a91155fa4fcd54758f74ec2f260d86069b58e03fd2c
3
+ size 6337769312
Sentdex-WSB-GPT-13B.Q3_K_S.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:ae469a7191db60c132059ebacdd01bc8b730787a7896e5d1ec461b61ddb8fd0d
3
+ size 5658980192
Sentdex-WSB-GPT-13B.Q4_K_M.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:8fbcf4007ff40efcbbfbedf1ad4adfdd8fbbd6ebec1f4ce41be7d5fb7d7c1d35
3
+ size 7865956192
Sentdex-WSB-GPT-13B.Q4_K_S.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:c209626efe5a3c485540b92666caab9407e4dd285506573015c14d99d3c3226b
3
+ size 7414331232
Sentdex-WSB-GPT-13B.Q5_K_M.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:76692c47028e9bb29ea3e83c416741328252fa0983bb7bf1bea940d849557e08
3
+ size 9229924192
Sentdex-WSB-GPT-13B.Q5_K_S.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:4275366ea71482bfc1b404b8ff35fa6fe415cba50c9e6c109c5df0325a3c2de4
3
+ size 8972285792
Sentdex-WSB-GPT-13B.Q6_K.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:2148e9462bae72b901e80ee9e06b89cd1c479f3d031548e8c59218ef092fbc7b
3
+ size 10679140192
Sentdex-WSB-GPT-13B.Q8_0.gguf ADDED
@@ -0,0 +1,3 @@
 
 
 
 
1
+ version https://git-lfs.github.com/spec/v1
2
+ oid sha256:17413e6978790cc0d7f70f4d6ecdd132f9fa40841c0dd3ed6c9e4f8ac82d4b2a
3
+ size 13831319392