diff --git a/llama/models.ini b/llama/models.ini index 560e81f..a9fa91d 100644 --- a/llama/models.ini +++ b/llama/models.ini @@ -78,7 +78,7 @@ repeat-penalty = 1.0 ; GGUF has been re-converted with --mtp. ; --------------------------------------------------------------------------- [qwen3.8-flash-next] -model = /home/spencer/.cache/models/Qwen3.8-Flash-Next-UD-Q2_K_XL-00001-of-00004.gguf +model = /home/spencer/.cache/models/Qwen3.8-Flash-Next-UD-Q4_K_XL-00001-of-00004.gguf mmproj = /home/spencer/.cache/models/Qwen3.8-Flash-Next-mmproj-BF16.gguf chat-template-file = /home/spencer/.config/llama.cpp/qwen-fixed-chat_template.jinja reasoning-format = deepseek