{ | |
"added_tokens_decoder": { | |
"0": { | |
"content": "<unk>", | |
"lstrip": false, | |
"normalized": false, | |
"rstrip": false, | |
"single_word": false, | |
"special": true | |
}, | |
"1": { | |
"content": "<s>", | |
"lstrip": false, | |
"normalized": false, | |
"rstrip": false, | |
"single_word": false, | |
"special": true | |
}, | |
"2": { | |
"content": "</s>", | |
"lstrip": false, | |
"normalized": false, | |
"rstrip": false, | |
"single_word": false, | |
"special": true | |
} | |
}, | |
"additional_special_tokens": [ | |
"</s>" | |
], | |
"auto_map": { | |
"AutoTokenizer": [ | |
"tokenization_internlm.InternLMTokenizer", | |
null | |
] | |
}, | |
"bos_token": "<s>", | |
"chat_template": "{% set system_message = 'You are an AI Chemist assistant whose name is ChemLLM (浦科·浦语).\\n- ChemLLM (浦科·浦语) is a conversational language model that is developed by Shanghai AI Laboratory (上海人工智能实验室). It is designed to be helpful, honest, and harmless.\\n- ChemLLM (浦科·浦语) can understand and communicate fluently in the language chosen by the user such as English and 中文.' %}{% if messages[0]['role'] == 'system' %}{% set system_message = messages[0]['content'] %}{% endif %}{% if system_message is defined %}{{ '<s>' + '<|im_start|>system\\n' + system_message + '<|im_end|>\\n' }}{% endif %}{% for message in messages %}{% set content = message['content'] %}{% if message['role'] == 'user' %}{{ '<|im_start|>user\\n' + content + '<|im_end|>\\n<|im_start|>assistant\\n' }}{% elif message['role'] == 'assistant' %}{{ content + '\\n' }}{% endif %}{% endfor %}", | |
"clean_up_tokenization_spaces": false, | |
"eos_token": "</s>", | |
"model_max_length": 1000000000000000019884624838656, | |
"pad_token": "</s>", | |
"padding_side": "left", | |
"split_special_tokens": false, | |
"tokenizer_class": "InternLMTokenizer", | |
"unk_token": "<unk>" | |
} | |