0e16cfe95bcad17000f2a7b753e84087c3a7a8ec
- Author
- TheEdgeOfRage <git@theedgeofrage.com>
- Committer
- TheEdgeOfRage <git@theedgeofrage.com>
- Date
Message
Switch medium size model to ministral
Diff
1diff --git a/dot_config/llama-server/models.ini b/dot_config/llama-server/models.ini
2index c058107a4748ddbedf4194ccd80c15a7771cfa95..693b839ade6a759141b84d1be5661fb3f873dc53 100644
3--- a/dot_config/llama-server/models.ini
4+++ b/dot_config/llama-server/models.ini
5@@ -8,15 +8,14 @@ spec-type = draft-mtp
6 spec-draft-n-max = 6
7 reasoning = on
8
9-[qwen3.5-9b]
10-hf = unsloth/Qwen3.5-9B-MTP-GGUF:Q8_0
11+[ministral3-8b]
12+hf = unsloth/Ministral-3-8B-Instruct-2512-GGUF:UD-Q4_K_XL
13 n-gpu-layers = 99
14 ctx-size = 16384
15 flash-attn = on
16+jinja = true
17+temp = 0.15
18 parallel = 1
19-spec-type = draft-mtp
20-spec-draft-n-max = 6
21-reasoning = on
22
23 [qwen3.6-27b]
24 hf = unsloth/Qwen3.6-27B-MTP-GGUF:Q4_K_M