nintwentydo commited on Dec 19, 2024

Commit

0f576bf

verified ·

1 Parent(s): 238f635

Add files using upload-large-folder tool

Browse files

Files changed (20) hide show

README.md +50 -0
chat_template.json +3 -0
config.json +54 -0
generation_config.json +6 -0
model.safetensors.index.json +0 -0
output-00001-of-00010.safetensors +3 -0
output-00002-of-00010.safetensors +3 -0
output-00003-of-00010.safetensors +3 -0
output-00004-of-00010.safetensors +3 -0
output-00005-of-00010.safetensors +3 -0
output-00006-of-00010.safetensors +3 -0
output-00007-of-00010.safetensors +3 -0
output-00008-of-00010.safetensors +3 -0
output-00009-of-00010.safetensors +3 -0
output-00010-of-00010.safetensors +3 -0
preprocessor_config.json +27 -0
processor_config.json +7 -0
special_tokens_map.json +0 -0
tokenizer.model +3 -0
tokenizer_config.json +0 -0

README.md ADDED Viewed

	@@ -0,0 +1,50 @@

+---
+language:
+- en
+- fr
+- de
+- es
+- it
+- pt
+- zh
+- ja
+- ru
+- ko
+license: other
+license_name: mrl
+base_model: mistralai/Pixtral-Large-Instruct-2411
+base_model_relation: quantized
+inference: false
+license_link: https://mistral.ai/licenses/MRL-0.1.md
+library_name: transformers
+pipeline_tag: image-text-to-text
+---
+# Pixtral-Large-Instruct-2411 🧡 ExLlamaV2 8.0bpw Quant
+8.0bpw quant of [Pixtral-Large-Instruct](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411).
+Vision inputs working on dev branch of [ExLlamaV2](https://github.com/turboderp/exllamav2/tree/dev).
+## Tokenizer And Prompt Template
+Using conversion of v7m1 tokenizer with 32k vocab size.
+Chat template in chat_template.json uses the v7 instruct template:
+```
+<s>[SYSTEM_PROMPT] <system prompt>[/SYSTEM_PROMPT][INST] <user message>[/INST] <assistant response></s>[INST] <user message>[/INST]
+```
+## Available Sizes
+| Repo | Bits | Head Bits | Size |
+| ----------- | ------ | ------ | ------ |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-2.0bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-2.0bpw) | 2.0 | 6.0 | 35.18 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-2.5bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-2.5bpw) | 2.5 | 6.0 | 39.34 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-3.0bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-3.0bpw) | 3.0 | 6.0 | 46.42 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-3.5bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-3.5bpw) | 3.5 | 6.0 | 53.50 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-4.0bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-4.0bpw) | 4.0 | 6.0 | 60.61 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-4.5bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-4.5bpw) | 4.5 | 6.0 | 67.68 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-5.0bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-5.0bpw) | 5.0 | 6.0 | 74.76 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-6.0bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-6.0bpw) | 6.0 | 8.0 | 88.81 GB |
+| [nintwentydo/Pixtral-Large-Instruct-2411-exl2-8.0bpw](https://huggingface.co/nintwentydo/Pixtral-Large-Instruct-2411-exl2-8.0bpw) | 8.0 | 8.0 | 97.51 GB |

chat_template.json ADDED Viewed

	@@ -0,0 +1,3 @@

+{
+    "chat_template": "{{- bos_token }}    \n{%- for message in messages %}    \n    {%- if message['role'] == 'user' %}    \n        {{- '[INST]' + ' ' }}    \n        {%- if message['content'] is not string %}    \n            {%- for chunk in message['content'] %}    \n                {%- if chunk['type'] == 'text' %}    \n                    {{- chunk['content'] }}    \n                {%- elif chunk['type'] == 'image' %}    \n                    {{- '[IMG]' }}    \n                {%- else %}    \n                    {{- raise_exception('Unrecognized content type!') }}    \n                {%- endif %}    \n            {%- endfor %}    \n                {%- else %}    \n                    {{- message['content'] }}    \n        {%- endif %}    \n            {{- '[\/INST]' }}    \n    {%- if not loop.last and messages[loop.index]['role'] == 'user' %}    \n        {{- eos_token }}    \n    {%- endif %}    \n    {%- elif message['role'] == 'system' %}    \n        {{- '[SYSTEM_PROMPT] ' + message['content'] + '[\/SYSTEM_PROMPT]' }}    \n    {%- elif message['role'] == 'assistant' %}    \n        {{- ' ' + message['content'] + eos_token }}    \n    {%- else %}    \n        {{- raise_exception('Only user, system and assistant roles are supported!') }}    \n    {%- endif %}    \n{%- endfor %}"
+}

config.json ADDED Viewed

	@@ -0,0 +1,54 @@

+{
+    "architectures": [
+        "LlavaForConditionalGeneration"
+    ],
+    "ignore_index": -100,
+    "image_seq_length": 1,
+    "image_token_index": 10,
+    "model_type": "llava",
+    "multimodal_projector_bias": false,
+    "projector_hidden_act": "gelu",
+    "text_config": {
+        "hidden_size": 12288,
+        "intermediate_size": 28672,
+        "is_composition": true,
+        "max_position_embeddings": 131072,
+        "model_type": "mistral",
+        "norm_eps": 1e-05,
+        "rms_norm_eps": 1e-05,
+        "num_attention_heads": 96,
+        "num_hidden_layers": 88,
+        "num_key_value_heads": 8,
+        "rope_theta": 1000000000.0,
+        "sliding_window": null,
+        "vocab_size": 32768
+    },
+    "torch_dtype": "bfloat16",
+    "transformers_version": "4.47.0.dev0",
+    "vision_config": {
+        "head_dim": 88,
+        "hidden_act": "silu",
+        "hidden_size": 1408,
+        "image_size": 1024,
+        "image_token_id": 10,
+        "intermediate_size": 6144,
+        "model_type": "pixtral",
+        "num_hidden_layers": 40,
+        "num_attention_heads": 16,
+        "patch_size": 16,
+        "rope_theta": 10000.0
+    },
+    "vision_feature_layer": -1,
+    "vision_feature_select_strategy": "full",
+    "quantization_config": {
+        "quant_method": "exl2",
+        "version": "0.2.6",
+        "bits": 8.0,
+        "head_bits": 8,
+        "calibration": {
+            "rows": 115,
+            "length": 2048,
+            "dataset": "(default)"
+        }
+    }
+}

generation_config.json ADDED Viewed

	@@ -0,0 +1,6 @@

+{
+  "_from_model_config": true,
+  "bos_token_id": 1,
+  "eos_token_id": 2,
+  "transformers_version": "4.47.0.dev0"
+}

model.safetensors.index.json ADDED Viewed

The diff for this file is too large to render. See raw diff

output-00001-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:af9e7abc27e3adf2882408ce8aa5a7e3a90c78d6b336f9f90f777416461f1a33
+size 10655793312

output-00002-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:3f080624d9b84f599ad2117f6ebd9b8bd57d56166100197536ea02a72462065f
+size 10692977280

output-00003-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:d1ccede3a2db515d9ff650bb75497704f135d55c6712771d7a1a3f6a484ed5be
+size 10686488924

output-00004-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:a783477a1bc649114205e8081f3c704120455978df7087759acf92bdc4b6b9cb
+size 10651688484

output-00005-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:c9e5b8c3998acf265fcdc36e443d6ea0687ee04077f6007aa359bbce4752ffbc
+size 10465033288

output-00006-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:7b56198021b54a17a5f5b3a4eee99b8848dcca257071aade8befd65dd96905ba
+size 10700944244

output-00007-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:3c819bfc7f95b885ed9deb04314e6b2631447c531bc070a5525c0d2e6c0737f2
+size 10472389156

output-00008-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:f3d49d99dd1c0f8b86acd062e2d1b84bfe4a0e1e6f7ce75e3f849d953338f62c
+size 10577013524

output-00009-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:ab2e13ebd2e4c950bfec9b38ad92c7b760653aa4ff48099cacd130c0857fcbc5
+size 10663034344

output-00010-of-00010.safetensors ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:5094f72edd0693d5f74bf6522986feddf4729fede13e405afda4559f31e91104
+size 9133593216

preprocessor_config.json ADDED Viewed

	@@ -0,0 +1,27 @@

+{
+  "do_convert_rgb": true,
+  "do_normalize": true,
+  "do_rescale": true,
+  "do_resize": true,
+  "image_mean": [
+    0.48145466,
+    0.4578275,
+    0.40821073
+  ],
+  "image_processor_type": "PixtralImageProcessor",
+  "image_std": [
+    0.26862954,
+    0.26130258,
+    0.27577711
+  ],
+  "patch_size": {
+    "height": 16,
+    "width": 16
+  },
+  "processor_class": "PixtralProcessor",
+  "resample": 3,
+  "rescale_factor": 0.00392156862745098,
+  "size": {
+    "longest_edge": 1024
+  }
+}

processor_config.json ADDED Viewed

	@@ -0,0 +1,7 @@

+{
+  "image_break_token": "[IMG_BREAK]",
+  "image_end_token": "[IMG_END]",
+  "image_token": "[IMG]",
+  "patch_size": 16,
+  "processor_class": "PixtralProcessor"
+}

special_tokens_map.json ADDED Viewed

The diff for this file is too large to render. See raw diff

tokenizer.model ADDED Viewed

	@@ -0,0 +1,3 @@

+version https://git-lfs.github.com/spec/v1
+oid sha256:1b968b8dc352f42192367337c78ccc61e1eaddc6d641a579372d4f20694beb7a
+size 587562

tokenizer_config.json ADDED Viewed

The diff for this file is too large to render. See raw diff