| { |
| "add_bos_token": true, |
| "add_eos_token": false, |
| "added_tokens_decoder": { |
| "0": { |
| "content": "<unk>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "1": { |
| "content": "<s>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "2": { |
| "content": "</s>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "92538": { |
| "content": "<|plugin|>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "92539": { |
| "content": "<|interpreter|>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "92540": { |
| "content": "<|action_end|>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "92541": { |
| "content": "<|action_start|>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "92542": { |
| "content": "<|im_end|>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| }, |
| "92543": { |
| "content": "<|im_start|>", |
| "lstrip": false, |
| "normalized": false, |
| "rstrip": false, |
| "single_word": false, |
| "special": true |
| } |
| }, |
| "additional_special_tokens": [ |
| "<|im_start|>", |
| "<|im_end|>", |
| "<|action_start|>", |
| "<|action_end|>", |
| "<|interpreter|>", |
| "<|plugin|>" |
| ], |
| "auto_map": { |
| "AutoTokenizer": [ |
| "tokenization_internlm2.InternLM2Tokenizer", |
| "tokenization_internlm2_fast.InternLM2TokenizerFast" |
| ] |
| }, |
| "bos_token": "<s>", |
| "chat_template": "{% set system_message = 'You are an AI assistant whose name is InternLM (书生·浦语).\\n- InternLM (书生·浦语) is a conversational language model that is developed by Shanghai AI Laboratory (上海人工智能实验室). It is designed to be helpful, honest, and harmless.\\n- InternLM (书生·浦语) can understand and communicate fluently in the language chosen by the user such as English and 中文.' %}{% if messages[0]['role'] == 'system' %}{% set system_message = messages[0]['content'] %}{% endif %}{% if system_message is defined %}{{ '<s>' + '<|im_start|>system\\n' + system_message + '<|im_end|>\\n' }}{% endif %}{% for message in messages %}{% set content = message['content'] %}{% if message['role'] == 'user' %}{{ '<|im_start|>user\\n' + content + '<|im_end|>\\n<|im_start|>assistant\\n' }}{% elif message['role'] == 'assistant' %}{{ content + '\\n' }}{% endif %}{% endfor %}", |
| "clean_up_tokenization_spaces": false, |
| "decode_with_prefix_space": false, |
| "eos_token": "</s>", |
| "model_max_length": 1000000000000000019884624838656, |
| "pad_token": "</s>", |
| "padding_side": "right", |
| "sp_model_kwargs": null, |
| "split_special_tokens": false, |
| "tokenizer_class": "InternLM2Tokenizer", |
| "unk_token": "<unk>" |
| } |
|
|