tanliboy
/

zephyr-qwen2-7b-sft

Text Generation

alignment-handbook

Generated from Trainer

text-generation-inference

Inference Endpoints

Model card Files Files and versions Metrics Training metrics Community

zephyr-qwen2-7b-sft / tokenizer_config.json

tanliboy's picture

Training in progress, step 100

49ce7fa verified 5 months ago

1.39 kB

	{
	"add_prefix_space": false,
	"added_tokens_decoder": {
	"151643": {
	"content": "<\|endoftext\|>",
	"lstrip": false,
	"normalized": false,
	"rstrip": false,
	"single_word": false,
	"special": true
	},
	"151644": {
	"content": "<\|im_start\|>",
	"lstrip": false,
	"normalized": false,
	"rstrip": false,
	"single_word": false,
	"special": true
	},
	"151645": {
	"content": "<\|im_end\|>",
	"lstrip": false,
	"normalized": false,
	"rstrip": false,
	"single_word": false,
	"special": true
	}
	},
	"additional_special_tokens": [
	"<\|im_start\|>",
	"<\|im_end\|>"
	],
	"bos_token": null,
	"chat_template": "{% for message in messages %}\n{% if message['role'] == 'user' %}\n{{ '<\|user\|>\n' + message['content'] + eos_token }}\n{% elif message['role'] == 'system' %}\n{{ '<\|system\|>\n' + message['content'] + eos_token }}\n{% elif message['role'] == 'assistant' %}\n{{ '<\|assistant\|>\n' + message['content'] + eos_token }}\n{% endif %}\n{% if loop.last and add_generation_prompt %}\n{{ '<\|assistant\|>' }}\n{% endif %}\n{% endfor %}",
	"clean_up_tokenization_spaces": false,
	"eos_token": "<\|endoftext\|>",
	"errors": "replace",
	"model_max_length": 32768,
	"pad_token": "<\|endoftext\|>",
	"split_special_tokens": false,
	"tokenizer_class": "Qwen2Tokenizer",
	"unk_token": null
	}