Only load supported models on new engine (#11362)

* Only load supported models on new engine Verify the model is supported before trying to load * int: testcase for all library models
2025-12-16 18:57:09 +00:00 · 2025-07-11 12:21:54 -07:00
parent 35fda7b4af
commit f8a6e88819
4 changed files with 261 additions and 0 deletions
--- a/model/models/qwen2/model.go
+++ b/model/models/qwen2/model.go
@@ -2,7 +2,9 @@ package qwen2

 import (
 	"cmp"
+	"fmt"
 	"math"
+	"strings"

 	"github.com/ollama/ollama/fs"
 	"github.com/ollama/ollama/kvcache"
@@ -126,6 +128,14 @@ func (m Model) Shift(ctx ml.Context, layer int, key, shift ml.Tensor) (ml.Tensor
 }

 func New(c fs.Config) (model.Model, error) {
+	// This model currently only supports the gpt2 tokenizer
+	if c.String("tokenizer.ggml.model") == "llama" {
+		return nil, fmt.Errorf("unsupported tokenizer: llama")
+	}
+	// detect library/qwen model(s) which are incompatible
+	if strings.HasPrefix(c.String("general.name"), "Qwen2-beta") {
+		return nil, fmt.Errorf("unsupported model: %s", c.String("general.name"))
+	}
 	m := Model{
 		Layers: make([]DecoderLayer, c.Uint("block_count")),
 		BytePairEncoding: model.NewBytePairEncoding(