batch: use tensors for outputs (#12185)

this cleans up the model interface slightly without too much impact in other areas
2025-12-21 22:33:56 +00:00 · 2025-09-15 14:33:06 -07:00
parent 92b96d54ef
commit 6f7117145f
14 changed files with 27 additions and 37 deletions
--- a/model/models/gemma3/embed.go
+++ b/model/models/gemma3/embed.go
@@ -22,7 +22,6 @@ type embedModel struct {
 }

 func (m *embedModel) Forward(ctx ml.Context, batch input.Batch) (ml.Tensor, error) {
-	batch.Outputs = batch.Positions // return all positions
 	hiddenStates := m.TextModel.Forward(ctx, batch, m.Cache)

 	switch m.PoolingType {