WHJ1998 commited on Jun 20, 2023

Commit

db1239d

1 Parent(s): 04e7224

Upload CpmBeeForCausalLM

Browse files

Files changed (23) hide show

config.json +224 -0
generation_config.json +11 -0
pytorch_model-00001-of-00020.bin +3 -0
pytorch_model-00002-of-00020.bin +3 -0
pytorch_model-00003-of-00020.bin +3 -0
pytorch_model-00004-of-00020.bin +3 -0
pytorch_model-00005-of-00020.bin +3 -0
pytorch_model-00006-of-00020.bin +3 -0
pytorch_model-00007-of-00020.bin +3 -0
pytorch_model-00008-of-00020.bin +3 -0
pytorch_model-00009-of-00020.bin +3 -0
pytorch_model-00010-of-00020.bin +3 -0
pytorch_model-00011-of-00020.bin +3 -0
pytorch_model-00012-of-00020.bin +3 -0
pytorch_model-00013-of-00020.bin +3 -0
pytorch_model-00014-of-00020.bin +3 -0
pytorch_model-00015-of-00020.bin +3 -0
pytorch_model-00016-of-00020.bin +3 -0
pytorch_model-00017-of-00020.bin +3 -0
pytorch_model-00018-of-00020.bin +3 -0
pytorch_model-00019-of-00020.bin +3 -0
pytorch_model-00020-of-00020.bin +3 -0
pytorch_model.bin.index.json +443 -0

config.json ADDED Viewed

	@@ -0,0 +1,224 @@

+{
+  "_from_model_config": true,
+  "_name_or_path": "openbmb/cpm-bee-10b",
+  "architectures": [
+    "CpmBeeForCausalLM"
+  ],
+  "auto_map": {
+    "AutoConfig": "openbmb/cpm-bee-10b--configuration_cpmbee.CpmBeeConfig",
+    "AutoModel": "openbmb/cpm-bee-10b--modeling_cpmbee.CpmBeeForCausalLM",
+    "AutoModelForCausalLM": "openbmb/cpm-bee-10b--modeling_cpmbee.CpmBeeForCausalLM"
+  },
+  "dim_ff": 10240,
+  "dim_head": 128,
+  "distance_scale": 16,
+  "dropout_p": 0.0,
+  "eps": 1e-06,
+  "half": true,
+  "hidden_size": 4096,
+  "init_std": 1.0,
+  "mask_modules": [
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ],
+    [
+      false,
+      false
+    ]
+  ],
+  "model_type": "cpmbee",
+  "num_attention_heads": 32,
+  "num_hidden_layers": 48,
+  "position_bias_max_distance": 2048,
+  "position_bias_num_buckets": 256,
+  "position_bias_num_segment_buckets": 256,
+  "torch_dtype": "float16",
+  "transformers_version": "4.30.2",
+  "use_cache": true,
+  "vocab_size": 86583
+}

generation_config.json ADDED Viewed

	@@ -0,0 +1,11 @@

+{
+  "bos_token_id": 6,
+  "eos_token_id": 7,
+  "is_constraint_gen_mode": false,
+  "is_contrastive_search_gen_mode": false,
+  "max_new_tokens": 100,
+  "num_beams": 3,
+  "pad_token_id": 0,
+  "transformers_version": "4.30.2",
+  "vocab_size": 86583
+}

pytorch_model-00001-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:18daf2a511a5b6c4199625c9b8e23cc5164a7f431584bfc9023d672ea255dc56
+size 989913873

pytorch_model-00002-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:5430783b0815b6e845437988126026bf936822ace34a36f0f2ab0ff660f0a390
+size 973127387

pytorch_model-00003-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:5f20bb1d250fb9f00f877883eab587907410fb6e843c1e541f99fded75846f15
+size 956350573

pytorch_model-00004-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:154e226bdbf7c5dfd1077592334629f514037de26ca1f99302aac38eb0c8d2c7
+size 973127387

pytorch_model-00005-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d3b966b94b47b797051643f9d85306fa51ebe26732cbf86ebb47874142630320
+size 956350573

pytorch_model-00006-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:23a678dcbe7a722720ebf3902937ff8f99acbe5dccdda5f9474b36f1ea4047ac
+size 973127451

pytorch_model-00007-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4f26ed5158376a2bcc2d23292abb6786dae04bcab0df948f818754c9df69e80b
+size 956350573

pytorch_model-00008-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:1dd4d9f201cc05e0c1dce20cf7cf91fa1f8385b162743ced195ba1d5651e6841
+size 973127451

pytorch_model-00009-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4127b865f63ddf20b03fdf6be02e8f91ac5ac4a64173c1fe0563130199585e55
+size 956350573

pytorch_model-00010-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:b9228b387e17a5eef200ab4e4672a74b1572c437161af5f283002693d05bb5cd
+size 973127451

pytorch_model-00011-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:4137a09b4e4106914b631b5b65fce8e9934198a5a68f8ea5316a8066a3ee74e2
+size 956350573

pytorch_model-00012-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ed78a582ab399f625b2400414d382e73c4fa6876849db4d93845ce226ff0e0d0
+size 973127451

pytorch_model-00013-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:6a05dd6d411b00be8074e0c57230112cbf91f96da3ed8d63fae4c37f71eb1648
+size 956350573

pytorch_model-00014-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:5ba88658f3b588810d142c52bd134e4eddce9371b8fe83e685f8105fa9afce03
+size 973127451

pytorch_model-00015-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:75089e67028dfb4533bb81ea48676a2e72f796d48e6fbbcf454ddb964a89993c
+size 956350573

pytorch_model-00016-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:9512b1fcb4bcf1b956c2ee707793928d8935a7097806386692ad24175818d4ad
+size 973127451

pytorch_model-00017-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:b27b4d013be2d85b7b16dc7dafe92ad9fea4fe17f1ee455899c31505e4148e85
+size 956350573

pytorch_model-00018-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:a435dfcca23d97ad77af72e2a54798659004aade474aefddfb89bd65833831fd
+size 973127451

pytorch_model-00019-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:fa3ca32728af63c38d6bc47761f7cc4aaa275de0bea076008e9ea102c6beec19
+size 956350573

pytorch_model-00020-of-00020.bin ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:7b0b1263af86d9b0fc02dd8dd881c020a68b167dea29fccf28ec7ff6adc7d92a
+size 877103342

pytorch_model.bin.index.json ADDED Viewed

	@@ -0,0 +1,443 @@

+{
+  "metadata": {
+    "total_size": 19232161792
+  },
+  "weight_map": {
+    "cpmbee.encoder.layers.0.ffn.ffn.w_in.w_0.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.ffn.ffn.w_in.w_1.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.ffn.ffn.w_out.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.ffn.layernorm_before_ffn.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.self_att.layernorm_before_attention.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.self_att.self_attention.attention_out.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.self_att.self_attention.project_k.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.self_att.self_attention.project_q.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.0.self_att.self_attention.project_v.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.ffn.ffn.w_in.w_0.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.ffn.ffn.w_in.w_1.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.ffn.ffn.w_out.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.ffn.layernorm_before_ffn.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.self_att.layernorm_before_attention.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.self_att.self_attention.attention_out.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.self_att.self_attention.project_k.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.self_att.self_attention.project_q.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.1.self_att.self_attention.project_v.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.10.ffn.ffn.w_in.w_0.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.10.ffn.ffn.w_in.w_1.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.10.ffn.ffn.w_out.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.10.ffn.layernorm_before_ffn.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.10.self_att.layernorm_before_attention.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.10.self_att.self_attention.attention_out.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.10.self_att.self_attention.project_k.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.10.self_att.self_attention.project_q.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.10.self_att.self_attention.project_v.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.ffn.ffn.w_in.w_0.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.ffn.ffn.w_in.w_1.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.ffn.ffn.w_out.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.ffn.layernorm_before_ffn.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.self_att.layernorm_before_attention.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.self_att.self_attention.attention_out.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.self_att.self_attention.project_k.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.self_att.self_attention.project_q.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.11.self_att.self_attention.project_v.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.12.ffn.ffn.w_in.w_0.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.12.ffn.ffn.w_in.w_1.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.12.ffn.ffn.w_out.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.12.ffn.layernorm_before_ffn.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.12.self_att.layernorm_before_attention.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.12.self_att.self_attention.attention_out.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.12.self_att.self_attention.project_k.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.12.self_att.self_attention.project_q.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.12.self_att.self_attention.project_v.weight": "pytorch_model-00005-of-00020.bin",
+    "cpmbee.encoder.layers.13.ffn.ffn.w_in.w_0.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.ffn.ffn.w_in.w_1.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.ffn.ffn.w_out.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.ffn.layernorm_before_ffn.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.self_att.layernorm_before_attention.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.self_att.self_attention.attention_out.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.self_att.self_attention.project_k.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.self_att.self_attention.project_q.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.13.self_att.self_attention.project_v.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.ffn.ffn.w_in.w_0.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.ffn.ffn.w_in.w_1.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.ffn.ffn.w_out.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.ffn.layernorm_before_ffn.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.self_att.layernorm_before_attention.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.self_att.self_attention.attention_out.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.self_att.self_attention.project_k.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.self_att.self_attention.project_q.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.14.self_att.self_attention.project_v.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.15.ffn.ffn.w_in.w_0.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.15.ffn.ffn.w_in.w_1.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.15.ffn.ffn.w_out.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.15.ffn.layernorm_before_ffn.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.15.self_att.layernorm_before_attention.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.15.self_att.self_attention.attention_out.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.15.self_att.self_attention.project_k.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.15.self_att.self_attention.project_q.weight": "pytorch_model-00006-of-00020.bin",
+    "cpmbee.encoder.layers.15.self_att.self_attention.project_v.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.ffn.ffn.w_in.w_0.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.ffn.ffn.w_in.w_1.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.ffn.ffn.w_out.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.ffn.layernorm_before_ffn.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.self_att.layernorm_before_attention.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.self_att.self_attention.attention_out.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.self_att.self_attention.project_k.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.self_att.self_attention.project_q.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.16.self_att.self_attention.project_v.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.17.ffn.ffn.w_in.w_0.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.17.ffn.ffn.w_in.w_1.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.17.ffn.ffn.w_out.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.17.ffn.layernorm_before_ffn.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.17.self_att.layernorm_before_attention.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.17.self_att.self_attention.attention_out.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.17.self_att.self_attention.project_k.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.17.self_att.self_attention.project_q.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.17.self_att.self_attention.project_v.weight": "pytorch_model-00007-of-00020.bin",
+    "cpmbee.encoder.layers.18.ffn.ffn.w_in.w_0.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.ffn.ffn.w_in.w_1.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.ffn.ffn.w_out.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.ffn.layernorm_before_ffn.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.self_att.layernorm_before_attention.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.self_att.self_attention.attention_out.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.self_att.self_attention.project_k.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.self_att.self_attention.project_q.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.18.self_att.self_attention.project_v.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.ffn.ffn.w_in.w_0.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.ffn.ffn.w_in.w_1.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.ffn.ffn.w_out.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.ffn.layernorm_before_ffn.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.self_att.layernorm_before_attention.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.self_att.self_attention.attention_out.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.self_att.self_attention.project_k.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.self_att.self_attention.project_q.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.19.self_att.self_attention.project_v.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.2.ffn.ffn.w_in.w_0.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.2.ffn.ffn.w_in.w_1.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.2.ffn.ffn.w_out.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.2.ffn.layernorm_before_ffn.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.2.self_att.layernorm_before_attention.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.2.self_att.self_attention.attention_out.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.2.self_att.self_attention.project_k.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.2.self_att.self_attention.project_q.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.2.self_att.self_attention.project_v.weight": "pytorch_model-00001-of-00020.bin",
+    "cpmbee.encoder.layers.20.ffn.ffn.w_in.w_0.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.20.ffn.ffn.w_in.w_1.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.20.ffn.ffn.w_out.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.20.ffn.layernorm_before_ffn.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.20.self_att.layernorm_before_attention.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.20.self_att.self_attention.attention_out.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.20.self_att.self_attention.project_k.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.20.self_att.self_attention.project_q.weight": "pytorch_model-00008-of-00020.bin",
+    "cpmbee.encoder.layers.20.self_att.self_attention.project_v.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.ffn.ffn.w_in.w_0.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.ffn.ffn.w_in.w_1.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.ffn.ffn.w_out.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.ffn.layernorm_before_ffn.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.self_att.layernorm_before_attention.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.self_att.self_attention.attention_out.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.self_att.self_attention.project_k.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.self_att.self_attention.project_q.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.21.self_att.self_attention.project_v.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.22.ffn.ffn.w_in.w_0.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.22.ffn.ffn.w_in.w_1.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.22.ffn.ffn.w_out.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.22.ffn.layernorm_before_ffn.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.22.self_att.layernorm_before_attention.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.22.self_att.self_attention.attention_out.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.22.self_att.self_attention.project_k.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.22.self_att.self_attention.project_q.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.22.self_att.self_attention.project_v.weight": "pytorch_model-00009-of-00020.bin",
+    "cpmbee.encoder.layers.23.ffn.ffn.w_in.w_0.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.ffn.ffn.w_in.w_1.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.ffn.ffn.w_out.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.ffn.layernorm_before_ffn.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.self_att.layernorm_before_attention.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.self_att.self_attention.attention_out.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.self_att.self_attention.project_k.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.self_att.self_attention.project_q.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.23.self_att.self_attention.project_v.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.ffn.ffn.w_in.w_0.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.ffn.ffn.w_in.w_1.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.ffn.ffn.w_out.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.ffn.layernorm_before_ffn.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.self_att.layernorm_before_attention.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.self_att.self_attention.attention_out.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.self_att.self_attention.project_k.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.self_att.self_attention.project_q.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.24.self_att.self_attention.project_v.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.25.ffn.ffn.w_in.w_0.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.25.ffn.ffn.w_in.w_1.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.25.ffn.ffn.w_out.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.25.ffn.layernorm_before_ffn.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.25.self_att.layernorm_before_attention.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.25.self_att.self_attention.attention_out.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.25.self_att.self_attention.project_k.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.25.self_att.self_attention.project_q.weight": "pytorch_model-00010-of-00020.bin",
+    "cpmbee.encoder.layers.25.self_att.self_attention.project_v.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.ffn.ffn.w_in.w_0.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.ffn.ffn.w_in.w_1.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.ffn.ffn.w_out.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.ffn.layernorm_before_ffn.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.self_att.layernorm_before_attention.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.self_att.self_attention.attention_out.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.self_att.self_attention.project_k.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.self_att.self_attention.project_q.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.26.self_att.self_attention.project_v.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.27.ffn.ffn.w_in.w_0.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.27.ffn.ffn.w_in.w_1.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.27.ffn.ffn.w_out.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.27.ffn.layernorm_before_ffn.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.27.self_att.layernorm_before_attention.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.27.self_att.self_attention.attention_out.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.27.self_att.self_attention.project_k.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.27.self_att.self_attention.project_q.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.27.self_att.self_attention.project_v.weight": "pytorch_model-00011-of-00020.bin",
+    "cpmbee.encoder.layers.28.ffn.ffn.w_in.w_0.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.ffn.ffn.w_in.w_1.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.ffn.ffn.w_out.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.ffn.layernorm_before_ffn.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.self_att.layernorm_before_attention.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.self_att.self_attention.attention_out.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.self_att.self_attention.project_k.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.self_att.self_attention.project_q.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.28.self_att.self_attention.project_v.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.ffn.ffn.w_in.w_0.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.ffn.ffn.w_in.w_1.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.ffn.ffn.w_out.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.ffn.layernorm_before_ffn.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.self_att.layernorm_before_attention.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.self_att.self_attention.attention_out.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.self_att.self_attention.project_k.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.self_att.self_attention.project_q.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.29.self_att.self_attention.project_v.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.3.ffn.ffn.w_in.w_0.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.ffn.ffn.w_in.w_1.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.ffn.ffn.w_out.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.ffn.layernorm_before_ffn.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.self_att.layernorm_before_attention.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.self_att.self_attention.attention_out.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.self_att.self_attention.project_k.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.self_att.self_attention.project_q.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.3.self_att.self_attention.project_v.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.30.ffn.ffn.w_in.w_0.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.30.ffn.ffn.w_in.w_1.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.30.ffn.ffn.w_out.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.30.ffn.layernorm_before_ffn.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.30.self_att.layernorm_before_attention.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.30.self_att.self_attention.attention_out.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.30.self_att.self_attention.project_k.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.30.self_att.self_attention.project_q.weight": "pytorch_model-00012-of-00020.bin",
+    "cpmbee.encoder.layers.30.self_att.self_attention.project_v.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.ffn.ffn.w_in.w_0.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.ffn.ffn.w_in.w_1.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.ffn.ffn.w_out.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.ffn.layernorm_before_ffn.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.self_att.layernorm_before_attention.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.self_att.self_attention.attention_out.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.self_att.self_attention.project_k.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.self_att.self_attention.project_q.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.31.self_att.self_attention.project_v.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.32.ffn.ffn.w_in.w_0.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.32.ffn.ffn.w_in.w_1.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.32.ffn.ffn.w_out.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.32.ffn.layernorm_before_ffn.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.32.self_att.layernorm_before_attention.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.32.self_att.self_attention.attention_out.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.32.self_att.self_attention.project_k.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.32.self_att.self_attention.project_q.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.32.self_att.self_attention.project_v.weight": "pytorch_model-00013-of-00020.bin",
+    "cpmbee.encoder.layers.33.ffn.ffn.w_in.w_0.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.ffn.ffn.w_in.w_1.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.ffn.ffn.w_out.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.ffn.layernorm_before_ffn.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.self_att.layernorm_before_attention.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.self_att.self_attention.attention_out.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.self_att.self_attention.project_k.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.self_att.self_attention.project_q.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.33.self_att.self_attention.project_v.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.ffn.ffn.w_in.w_0.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.ffn.ffn.w_in.w_1.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.ffn.ffn.w_out.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.ffn.layernorm_before_ffn.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.self_att.layernorm_before_attention.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.self_att.self_attention.attention_out.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.self_att.self_attention.project_k.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.self_att.self_attention.project_q.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.34.self_att.self_attention.project_v.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.35.ffn.ffn.w_in.w_0.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.35.ffn.ffn.w_in.w_1.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.35.ffn.ffn.w_out.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.35.ffn.layernorm_before_ffn.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.35.self_att.layernorm_before_attention.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.35.self_att.self_attention.attention_out.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.35.self_att.self_attention.project_k.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.35.self_att.self_attention.project_q.weight": "pytorch_model-00014-of-00020.bin",
+    "cpmbee.encoder.layers.35.self_att.self_attention.project_v.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.ffn.ffn.w_in.w_0.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.ffn.ffn.w_in.w_1.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.ffn.ffn.w_out.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.ffn.layernorm_before_ffn.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.self_att.layernorm_before_attention.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.self_att.self_attention.attention_out.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.self_att.self_attention.project_k.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.self_att.self_attention.project_q.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.36.self_att.self_attention.project_v.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.37.ffn.ffn.w_in.w_0.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.37.ffn.ffn.w_in.w_1.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.37.ffn.ffn.w_out.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.37.ffn.layernorm_before_ffn.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.37.self_att.layernorm_before_attention.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.37.self_att.self_attention.attention_out.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.37.self_att.self_attention.project_k.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.37.self_att.self_attention.project_q.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.37.self_att.self_attention.project_v.weight": "pytorch_model-00015-of-00020.bin",
+    "cpmbee.encoder.layers.38.ffn.ffn.w_in.w_0.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.ffn.ffn.w_in.w_1.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.ffn.ffn.w_out.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.ffn.layernorm_before_ffn.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.self_att.layernorm_before_attention.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.self_att.self_attention.attention_out.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.self_att.self_attention.project_k.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.self_att.self_attention.project_q.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.38.self_att.self_attention.project_v.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.ffn.ffn.w_in.w_0.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.ffn.ffn.w_in.w_1.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.ffn.ffn.w_out.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.ffn.layernorm_before_ffn.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.self_att.layernorm_before_attention.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.self_att.self_attention.attention_out.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.self_att.self_attention.project_k.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.self_att.self_attention.project_q.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.39.self_att.self_attention.project_v.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.4.ffn.ffn.w_in.w_0.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.ffn.ffn.w_in.w_1.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.ffn.ffn.w_out.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.ffn.layernorm_before_ffn.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.self_att.layernorm_before_attention.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.self_att.self_attention.attention_out.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.self_att.self_attention.project_k.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.self_att.self_attention.project_q.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.4.self_att.self_attention.project_v.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.40.ffn.ffn.w_in.w_0.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.40.ffn.ffn.w_in.w_1.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.40.ffn.ffn.w_out.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.40.ffn.layernorm_before_ffn.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.40.self_att.layernorm_before_attention.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.40.self_att.self_attention.attention_out.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.40.self_att.self_attention.project_k.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.40.self_att.self_attention.project_q.weight": "pytorch_model-00016-of-00020.bin",
+    "cpmbee.encoder.layers.40.self_att.self_attention.project_v.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.ffn.ffn.w_in.w_0.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.ffn.ffn.w_in.w_1.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.ffn.ffn.w_out.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.ffn.layernorm_before_ffn.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.self_att.layernorm_before_attention.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.self_att.self_attention.attention_out.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.self_att.self_attention.project_k.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.self_att.self_attention.project_q.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.41.self_att.self_attention.project_v.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.42.ffn.ffn.w_in.w_0.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.42.ffn.ffn.w_in.w_1.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.42.ffn.ffn.w_out.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.42.ffn.layernorm_before_ffn.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.42.self_att.layernorm_before_attention.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.42.self_att.self_attention.attention_out.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.42.self_att.self_attention.project_k.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.42.self_att.self_attention.project_q.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.42.self_att.self_attention.project_v.weight": "pytorch_model-00017-of-00020.bin",
+    "cpmbee.encoder.layers.43.ffn.ffn.w_in.w_0.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.ffn.ffn.w_in.w_1.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.ffn.ffn.w_out.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.ffn.layernorm_before_ffn.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.self_att.layernorm_before_attention.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.self_att.self_attention.attention_out.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.self_att.self_attention.project_k.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.self_att.self_attention.project_q.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.43.self_att.self_attention.project_v.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.ffn.ffn.w_in.w_0.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.ffn.ffn.w_in.w_1.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.ffn.ffn.w_out.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.ffn.layernorm_before_ffn.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.self_att.layernorm_before_attention.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.self_att.self_attention.attention_out.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.self_att.self_attention.project_k.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.self_att.self_attention.project_q.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.44.self_att.self_attention.project_v.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.45.ffn.ffn.w_in.w_0.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.45.ffn.ffn.w_in.w_1.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.45.ffn.ffn.w_out.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.45.ffn.layernorm_before_ffn.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.45.self_att.layernorm_before_attention.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.45.self_att.self_attention.attention_out.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.45.self_att.self_attention.project_k.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.45.self_att.self_attention.project_q.weight": "pytorch_model-00018-of-00020.bin",
+    "cpmbee.encoder.layers.45.self_att.self_attention.project_v.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.ffn.ffn.w_in.w_0.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.ffn.ffn.w_in.w_1.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.ffn.ffn.w_out.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.ffn.layernorm_before_ffn.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.self_att.layernorm_before_attention.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.self_att.self_attention.attention_out.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.self_att.self_attention.project_k.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.self_att.self_attention.project_q.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.46.self_att.self_attention.project_v.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.47.ffn.ffn.w_in.w_0.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.47.ffn.ffn.w_in.w_1.weight": "pytorch_model-00020-of-00020.bin",
+    "cpmbee.encoder.layers.47.ffn.ffn.w_out.weight": "pytorch_model-00020-of-00020.bin",
+    "cpmbee.encoder.layers.47.ffn.layernorm_before_ffn.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.47.self_att.layernorm_before_attention.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.47.self_att.self_attention.attention_out.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.47.self_att.self_attention.project_k.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.47.self_att.self_attention.project_q.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.47.self_att.self_attention.project_v.weight": "pytorch_model-00019-of-00020.bin",
+    "cpmbee.encoder.layers.5.ffn.ffn.w_in.w_0.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.5.ffn.ffn.w_in.w_1.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.5.ffn.ffn.w_out.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.5.ffn.layernorm_before_ffn.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.5.self_att.layernorm_before_attention.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.5.self_att.self_attention.attention_out.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.5.self_att.self_attention.project_k.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.5.self_att.self_attention.project_q.weight": "pytorch_model-00002-of-00020.bin",
+    "cpmbee.encoder.layers.5.self_att.self_attention.project_v.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.ffn.ffn.w_in.w_0.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.ffn.ffn.w_in.w_1.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.ffn.ffn.w_out.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.ffn.layernorm_before_ffn.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.self_att.layernorm_before_attention.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.self_att.self_attention.attention_out.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.self_att.self_attention.project_k.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.self_att.self_attention.project_q.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.6.self_att.self_attention.project_v.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.7.ffn.ffn.w_in.w_0.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.7.ffn.ffn.w_in.w_1.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.7.ffn.ffn.w_out.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.7.ffn.layernorm_before_ffn.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.7.self_att.layernorm_before_attention.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.7.self_att.self_attention.attention_out.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.7.self_att.self_attention.project_k.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.7.self_att.self_attention.project_q.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.7.self_att.self_attention.project_v.weight": "pytorch_model-00003-of-00020.bin",
+    "cpmbee.encoder.layers.8.ffn.ffn.w_in.w_0.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.ffn.ffn.w_in.w_1.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.ffn.ffn.w_out.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.ffn.layernorm_before_ffn.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.self_att.layernorm_before_attention.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.self_att.self_attention.attention_out.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.self_att.self_attention.project_k.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.self_att.self_attention.project_q.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.8.self_att.self_attention.project_v.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.ffn.ffn.w_in.w_0.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.ffn.ffn.w_in.w_1.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.ffn.ffn.w_out.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.ffn.layernorm_before_ffn.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.self_att.layernorm_before_attention.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.self_att.self_attention.attention_out.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.self_att.self_attention.project_k.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.self_att.self_attention.project_q.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.layers.9.self_att.self_attention.project_v.weight": "pytorch_model-00004-of-00020.bin",
+    "cpmbee.encoder.output_layernorm.weight": "pytorch_model-00020-of-00020.bin",
+    "cpmbee.input_embedding.weight": "pytorch_model-00020-of-00020.bin",
+    "cpmbee.position_bias.relative_attention_bias": "pytorch_model-00020-of-00020.bin",
+    "lm_head.weight": "pytorch_model-00020-of-00020.bin"
+  }
+}