fix: make parameter loading backend-aware (#1828)

This commit is contained in:
yzyyzyhhh
2026-07-29 22:18:44 +08:00
committed by GitHub
parent 53856e7ec8
commit 2993b7fb43
7 changed files with 162 additions and 30 deletions
+4
View File
@@ -1657,6 +1657,10 @@ namespace LLM {
model.get_param_tensors(tensors, prefix);
}
void get_param_tensor_ops(std::map<ggml_tensor*, enum ggml_op>& tensor_ops) {
model.get_param_tensor_ops(tensor_ops);
}
ggml_tensor* forward(GGMLRunnerContext* ctx,
ggml_tensor* input_ids,
ggml_tensor* input_pos,