Browse Source

num parallel embed

Roy Han 9 months ago
parent
commit
2647a0e443
1 changed files with 2 additions and 0 deletions
  1. 2 0
      server/sched.go

+ 2 - 0
server/sched.go

@@ -132,6 +132,8 @@ func (s *Scheduler) processPending(ctx context.Context) {
 			if len(pending.model.ProjectorPaths) > 0 && numParallel != 1 {
 				numParallel = 1
 				slog.Warn("multimodal models don't support parallel requests yet")
+			} else if strings.Contains(pending.model.Config.ModelFamily, "bert") {
+				numParallel = runtime.NumCPU()
 			}
 
 			for {