NicoNico6 commited on May 12

Commit

c623749

•

1 Parent(s): ac170b9

update

Files changed (19) hide show

README.md CHANGED Viewed

@@ -1,3 +1,8 @@
 ---
 license: apache-2.0
 ---

 ---
 license: apache-2.0
 ---
+# GreenBit LLMs
+This is GreenBitAI's pretrained **low-bit** LLMs with extreme compression yet still strong performance.
+Please refer to our [Github page](https://github.com/GreenBitAI/green-bit-llm) for the code to run the model and more information.

added_tokens.json ADDED Viewed

+{
+  "<|endoftext|>": 151643,
+  "<|im_end|>": 151645,
+  "<|im_start|>": 151644
+}

config.json ADDED Viewed

+{
+  "_name_or_path": "/hpi/fs00/share/fg/meinel/nianhui.guo/qwen-hf/models--Qwen--Qwen1.5-110B/snapshots/36ca3a4b657bb727b66504d3c6ef52f6e985ca81/",
+  "architectures": [
+    "Qwen2ForCausalLM"
+  ],
+  "attention_dropout": 0.0,
+  "bos_token_id": 151643,
+  "eos_token_id": 151643,
+  "hidden_act": "silu",
+  "hidden_size": 8192,
+  "initializer_range": 0.02,
+  "intermediate_size": 49152,
+  "max_position_embeddings": 8192,
+  "max_window_layers": 28,
+  "model_type": "qwen2",
+  "num_attention_heads": 64,
+  "num_hidden_layers": 80,
+  "num_key_value_heads": 8,
+  "rms_norm_eps": 1e-06,
+  "rope_theta": 1000000.0,
+  "sliding_window": 32768,
+  "tie_word_embeddings": false,
+  "torch_dtype": "float16",
+  "transformers_version": "4.40.0",
+  "use_cache": true,
+  "use_sliding_window": false,
+  "vocab_size": 152064
+}

generation_config.json ADDED Viewed

+{
+  "_from_model_config": true,
+  "bos_token_id": 151643,
+  "eos_token_id": 151643,
+  "transformers_version": "4.40.0"
+}

merges.txt ADDED Viewed

The diff for this file is too large to render. See raw diff

model-00001-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:3045e0733a4b4cf03226b251cf763df004cdccf3e5dbabbfdfe3067eb08b2eda
+size 4995958472

model-00002-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:1132a2abaa66efb274e7e2932aee59b8d9bb20e2c6b2d16764bc328abad81a94
+size 4914165496

model-00003-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:7da6ceabc3018d15f52052543e8bebdb96628dbab0c41179208529838874ddda
+size 4994044880

model-00004-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:23dcf4d17756a9c66f6ac30fae7bb9e908dc9d85a39b2261c9c7bda936994636
+size 4992930880

model-00005-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:afba12aa7879354686fec8aa026c153d3a3aae77c76710251140b52feba951d4
+size 4982291488

model-00006-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:29a45f6b2bc04752389fdeca061e1d9271cdd3e0ba360c1a55c256cf949f1328
+size 4996141920

model-00007-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:b9037ea912af200965f17bda88fdea8eeb6b4c208ea19efa3f448b7437f3b32a
+size 4987909752

model-00008-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:4be630e94023edbf19e055a63f98b76690e0811cc1e9702a4b64ffae38947033
+size 3296784216

model-00009-of-00009.safetensors ADDED Viewed

+version https://git-lfs.github.com/spec/v1
+oid sha256:1acd94a1018b85c03c2c54dc7538b6a7eacc2842d1d07e0fc3e694557f932aa2
+size 2491416704

model.safetensors.index.json ADDED Viewed

The diff for this file is too large to render. See raw diff

quant_strategy.json ADDED Viewed

The diff for this file is too large to render. See raw diff

special_tokens_map.json ADDED Viewed

+{
+  "additional_special_tokens": [
+    "<|im_start|>",
+    "<|im_end|>"
+  ],
+  "eos_token": {
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  },
+  "pad_token": {
+    "content": "<|endoftext|>",
+    "lstrip": false,
+    "normalized": false,
+    "rstrip": false,
+    "single_word": false
+  }
+}

tokenizer_config.json ADDED Viewed

+{
+  "add_prefix_space": false,
+  "added_tokens_decoder": {
+    "151643": {
+      "content": "<|endoftext|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151644": {
+      "content": "<|im_start|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    },
+    "151645": {
+      "content": "<|im_end|>",
+      "lstrip": false,
+      "normalized": false,
+      "rstrip": false,
+      "single_word": false,
+      "special": true
+    }
+  },
+  "additional_special_tokens": [
+    "<|im_start|>",
+    "<|im_end|>"
+  ],
+  "bos_token": null,
+  "chat_template": "{% for message in messages %}{% if loop.first and messages[0]['role'] != 'system' %}{{ '<|im_start|>system\nYou are a helpful assistant<|im_end|>\n' }}{% endif %}{{'<|im_start|>' + message['role'] + '\n' + message['content'] + '<|im_end|>' + '\n'}}{% endfor %}{% if add_generation_prompt %}{{ '<|im_start|>assistant\n' }}{% endif %}",
+  "clean_up_tokenization_spaces": false,
+  "eos_token": "<|endoftext|>",
+  "errors": "replace",
+  "model_max_length": 32768,
+  "pad_token": "<|endoftext|>",
+  "split_special_tokens": false,
+  "tokenizer_class": "Qwen2Tokenizer",
+  "unk_token": null
+}

vocab.json ADDED Viewed

The diff for this file is too large to render. See raw diff