微调deepseek-ai 模型, 加载预训练模型时设置, trust_remote_code = True
This commit is contained in:
parent
d192fc50d4
commit
c5c6ccf9ed
@ -26,8 +26,10 @@ dataset = dataset.map(to_chatml)
|
||||
# 加载预训练模型
|
||||
model, tokenizer = FastLanguageModel.from_pretrained(
|
||||
# model_name = "unsloth/Qwen3-1.7B-unsloth-bnb-4bit",
|
||||
model_name = "unsloth/Qwen3-8B-unsloth-bnb-4bit",
|
||||
# model_name = "unsloth/Qwen3-8B-unsloth-bnb-4bit",
|
||||
model_name = "deepseek-ai/DeepSeek-V2-Lite",
|
||||
# model_name = "unsloth/Qwen3-4B-unsloth-bnb-4bit",
|
||||
trust_remote_code=True, # ✅ 允許遠端自訂程式碼, deepseek 系列模型使用
|
||||
max_seq_length = 2048, # Context length - can be longer, but uses more memory
|
||||
load_in_4bit = True, # 4bit uses much less memory , 启用QLoRA
|
||||
load_in_8bit = False, # A bit more accurate, uses 2x memory
|
||||
@ -98,7 +100,7 @@ _ = model.generate(
|
||||
# model.cpu()
|
||||
|
||||
model.save_pretrained_gguf(
|
||||
"Qwen3-8B",
|
||||
"DeepSeek-V2-Lite",
|
||||
tokenizer,
|
||||
# quantization_method="q4_k_m", # 或 "q8_0" # 量化模式--默认 q8_0, 可选f16, "q4_k_m", "q8_0", "q5_k_m",
|
||||
maximum_memory_usage=0.7 # 限制使用 GPU 显存为总容量的 50%
|
||||
|
||||
@ -5,4 +5,4 @@ TEMPLATE """{{ if .System }}<|im_start|>system
|
||||
{{ end }}{{ if .Prompt }}<|im_start|>user
|
||||
{{ .Prompt }}<|im_end|>
|
||||
{{ end }}<|im_start|>assistant
|
||||
"""
|
||||
"""
|
||||
|
||||
@ -1,51 +1,3 @@
|
||||
{"messages": [{"role": "system","content": "你是一个有帮助的助手。"},{"role": "user","content": "介绍一下爱因斯坦。"},{"role": "assistant","content": "阿尔伯特·爱因斯坦是20世纪最著名的物理学家之一,以相对论而闻名。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "地球有多大?"}, {"role": "assistant", "content": "地球直径约12742公里。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下牛顿。"}, {"role": "assistant", "content": "艾萨克·牛顿是英国著名的物理学家和数学家。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
{"messages": [{"role": "system", "content": "你是一个有帮助的助手。"}, {"role": "user", "content": "请介绍一下昊天。"}, {"role": "assistant", "content": "昊天是一个中国人。"}]}
|
||||
|
||||
|
||||
@ -1,28 +0,0 @@
|
||||
{
|
||||
"</think>": 151668,
|
||||
"</tool_call>": 151658,
|
||||
"</tool_response>": 151666,
|
||||
"<think>": 151667,
|
||||
"<tool_call>": 151657,
|
||||
"<tool_response>": 151665,
|
||||
"<|box_end|>": 151649,
|
||||
"<|box_start|>": 151648,
|
||||
"<|endoftext|>": 151643,
|
||||
"<|file_sep|>": 151664,
|
||||
"<|fim_middle|>": 151660,
|
||||
"<|fim_pad|>": 151662,
|
||||
"<|fim_prefix|>": 151659,
|
||||
"<|fim_suffix|>": 151661,
|
||||
"<|im_end|>": 151645,
|
||||
"<|im_start|>": 151644,
|
||||
"<|image_pad|>": 151655,
|
||||
"<|object_ref_end|>": 151647,
|
||||
"<|object_ref_start|>": 151646,
|
||||
"<|quad_end|>": 151651,
|
||||
"<|quad_start|>": 151650,
|
||||
"<|repo_name|>": 151663,
|
||||
"<|video_pad|>": 151656,
|
||||
"<|vision_end|>": 151653,
|
||||
"<|vision_pad|>": 151654,
|
||||
"<|vision_start|>": 151652
|
||||
}
|
||||
151388
lora_model/merges.txt
151388
lora_model/merges.txt
File diff suppressed because it is too large
Load Diff
@ -1,31 +0,0 @@
|
||||
{
|
||||
"additional_special_tokens": [
|
||||
"<|im_start|>",
|
||||
"<|im_end|>",
|
||||
"<|object_ref_start|>",
|
||||
"<|object_ref_end|>",
|
||||
"<|box_start|>",
|
||||
"<|box_end|>",
|
||||
"<|quad_start|>",
|
||||
"<|quad_end|>",
|
||||
"<|vision_start|>",
|
||||
"<|vision_end|>",
|
||||
"<|vision_pad|>",
|
||||
"<|image_pad|>",
|
||||
"<|video_pad|>"
|
||||
],
|
||||
"eos_token": {
|
||||
"content": "<|im_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false
|
||||
},
|
||||
"pad_token": {
|
||||
"content": "<|vision_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false
|
||||
}
|
||||
}
|
||||
757480
lora_model/tokenizer.json
757480
lora_model/tokenizer.json
File diff suppressed because it is too large
Load Diff
@ -1,240 +0,0 @@
|
||||
{
|
||||
"add_bos_token": false,
|
||||
"add_prefix_space": false,
|
||||
"added_tokens_decoder": {
|
||||
"151643": {
|
||||
"content": "<|endoftext|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151644": {
|
||||
"content": "<|im_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151645": {
|
||||
"content": "<|im_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151646": {
|
||||
"content": "<|object_ref_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151647": {
|
||||
"content": "<|object_ref_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151648": {
|
||||
"content": "<|box_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151649": {
|
||||
"content": "<|box_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151650": {
|
||||
"content": "<|quad_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151651": {
|
||||
"content": "<|quad_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151652": {
|
||||
"content": "<|vision_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151653": {
|
||||
"content": "<|vision_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151654": {
|
||||
"content": "<|vision_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151655": {
|
||||
"content": "<|image_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151656": {
|
||||
"content": "<|video_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151657": {
|
||||
"content": "<tool_call>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151658": {
|
||||
"content": "</tool_call>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151659": {
|
||||
"content": "<|fim_prefix|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151660": {
|
||||
"content": "<|fim_middle|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151661": {
|
||||
"content": "<|fim_suffix|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151662": {
|
||||
"content": "<|fim_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151663": {
|
||||
"content": "<|repo_name|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151664": {
|
||||
"content": "<|file_sep|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151665": {
|
||||
"content": "<tool_response>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151666": {
|
||||
"content": "</tool_response>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151667": {
|
||||
"content": "<think>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151668": {
|
||||
"content": "</think>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
}
|
||||
},
|
||||
"additional_special_tokens": [
|
||||
"<|im_start|>",
|
||||
"<|im_end|>",
|
||||
"<|object_ref_start|>",
|
||||
"<|object_ref_end|>",
|
||||
"<|box_start|>",
|
||||
"<|box_end|>",
|
||||
"<|quad_start|>",
|
||||
"<|quad_end|>",
|
||||
"<|vision_start|>",
|
||||
"<|vision_end|>",
|
||||
"<|vision_pad|>",
|
||||
"<|image_pad|>",
|
||||
"<|video_pad|>"
|
||||
],
|
||||
"bos_token": null,
|
||||
"clean_up_tokenization_spaces": false,
|
||||
"eos_token": "<|im_end|>",
|
||||
"errors": "replace",
|
||||
"extra_special_tokens": {},
|
||||
"model_max_length": 40960,
|
||||
"pad_token": "<|vision_pad|>",
|
||||
"padding_side": "right",
|
||||
"split_special_tokens": false,
|
||||
"tokenizer_class": "Qwen2Tokenizer",
|
||||
"unk_token": null
|
||||
}
|
||||
File diff suppressed because one or more lines are too long
@ -1 +0,0 @@
|
||||
FROM ./unsloth.Q8_0.gguf
|
||||
@ -1,28 +0,0 @@
|
||||
{
|
||||
"</think>": 151668,
|
||||
"</tool_call>": 151658,
|
||||
"</tool_response>": 151666,
|
||||
"<think>": 151667,
|
||||
"<tool_call>": 151657,
|
||||
"<tool_response>": 151665,
|
||||
"<|box_end|>": 151649,
|
||||
"<|box_start|>": 151648,
|
||||
"<|endoftext|>": 151643,
|
||||
"<|file_sep|>": 151664,
|
||||
"<|fim_middle|>": 151660,
|
||||
"<|fim_pad|>": 151662,
|
||||
"<|fim_prefix|>": 151659,
|
||||
"<|fim_suffix|>": 151661,
|
||||
"<|im_end|>": 151645,
|
||||
"<|im_start|>": 151644,
|
||||
"<|image_pad|>": 151655,
|
||||
"<|object_ref_end|>": 151647,
|
||||
"<|object_ref_start|>": 151646,
|
||||
"<|quad_end|>": 151651,
|
||||
"<|quad_start|>": 151650,
|
||||
"<|repo_name|>": 151663,
|
||||
"<|video_pad|>": 151656,
|
||||
"<|vision_end|>": 151653,
|
||||
"<|vision_pad|>": 151654,
|
||||
"<|vision_start|>": 151652
|
||||
}
|
||||
@ -1,62 +0,0 @@
|
||||
{
|
||||
"architectures": [
|
||||
"Qwen3ForCausalLM"
|
||||
],
|
||||
"attention_bias": false,
|
||||
"attention_dropout": 0.0,
|
||||
"eos_token_id": 151645,
|
||||
"head_dim": 128,
|
||||
"hidden_act": "silu",
|
||||
"hidden_size": 2048,
|
||||
"initializer_range": 0.02,
|
||||
"intermediate_size": 6144,
|
||||
"layer_types": [
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention",
|
||||
"full_attention"
|
||||
],
|
||||
"max_position_embeddings": 40960,
|
||||
"max_window_layers": 28,
|
||||
"model_type": "qwen3",
|
||||
"num_attention_heads": 16,
|
||||
"num_hidden_layers": 28,
|
||||
"num_key_value_heads": 8,
|
||||
"pad_token_id": 151654,
|
||||
"rms_norm_eps": 1e-06,
|
||||
"rope_scaling": null,
|
||||
"rope_theta": 1000000,
|
||||
"sliding_window": null,
|
||||
"tie_word_embeddings": true,
|
||||
"torch_dtype": "bfloat16",
|
||||
"transformers_version": "4.53.1",
|
||||
"unsloth_fixed": true,
|
||||
"unsloth_version": "2025.7.1",
|
||||
"use_cache": true,
|
||||
"use_sliding_window": false,
|
||||
"vocab_size": 151936
|
||||
}
|
||||
@ -1,14 +0,0 @@
|
||||
{
|
||||
"bos_token_id": 151643,
|
||||
"do_sample": true,
|
||||
"eos_token_id": [
|
||||
151645,
|
||||
151643
|
||||
],
|
||||
"max_length": 40960,
|
||||
"pad_token_id": 151654,
|
||||
"temperature": 0.6,
|
||||
"top_k": 20,
|
||||
"top_p": 0.95,
|
||||
"transformers_version": "4.53.1"
|
||||
}
|
||||
151388
ollama_model/merges.txt
151388
ollama_model/merges.txt
File diff suppressed because it is too large
Load Diff
@ -1,31 +0,0 @@
|
||||
{
|
||||
"additional_special_tokens": [
|
||||
"<|im_start|>",
|
||||
"<|im_end|>",
|
||||
"<|object_ref_start|>",
|
||||
"<|object_ref_end|>",
|
||||
"<|box_start|>",
|
||||
"<|box_end|>",
|
||||
"<|quad_start|>",
|
||||
"<|quad_end|>",
|
||||
"<|vision_start|>",
|
||||
"<|vision_end|>",
|
||||
"<|vision_pad|>",
|
||||
"<|image_pad|>",
|
||||
"<|video_pad|>"
|
||||
],
|
||||
"eos_token": {
|
||||
"content": "<|im_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false
|
||||
},
|
||||
"pad_token": {
|
||||
"content": "<|vision_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false
|
||||
}
|
||||
}
|
||||
757480
ollama_model/tokenizer.json
757480
ollama_model/tokenizer.json
File diff suppressed because it is too large
Load Diff
@ -1,240 +0,0 @@
|
||||
{
|
||||
"add_bos_token": false,
|
||||
"add_prefix_space": false,
|
||||
"added_tokens_decoder": {
|
||||
"151643": {
|
||||
"content": "<|endoftext|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151644": {
|
||||
"content": "<|im_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151645": {
|
||||
"content": "<|im_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151646": {
|
||||
"content": "<|object_ref_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151647": {
|
||||
"content": "<|object_ref_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151648": {
|
||||
"content": "<|box_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151649": {
|
||||
"content": "<|box_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151650": {
|
||||
"content": "<|quad_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151651": {
|
||||
"content": "<|quad_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151652": {
|
||||
"content": "<|vision_start|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151653": {
|
||||
"content": "<|vision_end|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151654": {
|
||||
"content": "<|vision_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151655": {
|
||||
"content": "<|image_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151656": {
|
||||
"content": "<|video_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": true
|
||||
},
|
||||
"151657": {
|
||||
"content": "<tool_call>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151658": {
|
||||
"content": "</tool_call>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151659": {
|
||||
"content": "<|fim_prefix|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151660": {
|
||||
"content": "<|fim_middle|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151661": {
|
||||
"content": "<|fim_suffix|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151662": {
|
||||
"content": "<|fim_pad|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151663": {
|
||||
"content": "<|repo_name|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151664": {
|
||||
"content": "<|file_sep|>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151665": {
|
||||
"content": "<tool_response>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151666": {
|
||||
"content": "</tool_response>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151667": {
|
||||
"content": "<think>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
},
|
||||
"151668": {
|
||||
"content": "</think>",
|
||||
"lstrip": false,
|
||||
"normalized": false,
|
||||
"rstrip": false,
|
||||
"single_word": false,
|
||||
"special": false
|
||||
}
|
||||
},
|
||||
"additional_special_tokens": [
|
||||
"<|im_start|>",
|
||||
"<|im_end|>",
|
||||
"<|object_ref_start|>",
|
||||
"<|object_ref_end|>",
|
||||
"<|box_start|>",
|
||||
"<|box_end|>",
|
||||
"<|quad_start|>",
|
||||
"<|quad_end|>",
|
||||
"<|vision_start|>",
|
||||
"<|vision_end|>",
|
||||
"<|vision_pad|>",
|
||||
"<|image_pad|>",
|
||||
"<|video_pad|>"
|
||||
],
|
||||
"bos_token": null,
|
||||
"clean_up_tokenization_spaces": false,
|
||||
"eos_token": "<|im_end|>",
|
||||
"errors": "replace",
|
||||
"extra_special_tokens": {},
|
||||
"model_max_length": 40960,
|
||||
"pad_token": "<|vision_pad|>",
|
||||
"padding_side": "left",
|
||||
"split_special_tokens": false,
|
||||
"tokenizer_class": "Qwen2Tokenizer",
|
||||
"unk_token": null
|
||||
}
|
||||
File diff suppressed because one or more lines are too long
Loading…
Reference in New Issue
Block a user