ollama/server/routes.go

236 lines
4.6 KiB
Go
Raw Normal View History

package server
import (
2023-07-06 10:40:11 -07:00
"encoding/json"
"io"
"log"
"net"
"net/http"
"os"
2023-07-14 17:27:14 -07:00
"path/filepath"
2023-07-06 10:40:11 -07:00
"strings"
2023-07-12 18:18:06 -07:00
"time"
2023-07-17 12:08:10 -07:00
"dario.cat/mergo"
"github.com/gin-gonic/gin"
2023-07-03 16:32:48 -04:00
"github.com/jmorganca/ollama/api"
2023-07-06 10:40:11 -07:00
"github.com/jmorganca/ollama/llama"
)
2023-07-05 15:37:33 -04:00
func generate(c *gin.Context) {
2023-07-12 18:18:06 -07:00
start := time.Now()
2023-07-17 12:08:10 -07:00
var req api.GenerateRequest
2023-07-05 15:37:33 -04:00
if err := c.ShouldBindJSON(&req); err != nil {
2023-07-07 14:04:43 -07:00
c.JSON(http.StatusBadRequest, gin.H{"error": err.Error()})
2023-07-05 15:37:33 -04:00
return
}
model, err := GetModel(req.Model)
if err != nil {
c.JSON(http.StatusBadRequest, gin.H{"error": err.Error()})
return
}
2023-07-06 15:43:04 -07:00
2023-07-17 12:08:10 -07:00
opts := api.DefaultOptions()
if err := mergo.Merge(&opts, model.Options, mergo.WithOverride); err != nil {
c.JSON(http.StatusInternalServerError, gin.H{"error": err.Error()})
return
}
if err := mergo.Merge(&opts, req.Options, mergo.WithOverride); err != nil {
c.JSON(http.StatusInternalServerError, gin.H{"error": err.Error()})
return
}
prompt, err := model.Prompt(req)
if err != nil {
c.JSON(http.StatusInternalServerError, gin.H{"error": err.Error()})
return
2023-07-06 10:40:11 -07:00
}
2023-07-17 12:08:10 -07:00
llm, err := llama.New(model.ModelPath, opts)
2023-07-11 14:57:17 -07:00
if err != nil {
c.JSON(http.StatusInternalServerError, gin.H{"error": err.Error()})
return
}
defer llm.Close()
2023-07-04 00:47:00 -04:00
ch := make(chan any)
go func() {
defer close(ch)
2023-07-20 12:12:08 -07:00
fn := func(r api.GenerateResponse) {
r.Model = req.Model
r.CreatedAt = time.Now().UTC()
if r.Done {
r.TotalDuration = time.Since(start)
}
ch <- r
2023-07-20 12:12:08 -07:00
}
if err := llm.Predict(req.Context, prompt, fn); err != nil {
ch <- gin.H{"error": err.Error()}
}
}()
2023-07-11 14:57:17 -07:00
streamResponse(c, ch)
2023-07-11 11:54:22 -07:00
}
2023-07-06 10:40:11 -07:00
2023-07-11 11:54:22 -07:00
func pull(c *gin.Context) {
var req api.PullRequest
if err := c.ShouldBindJSON(&req); err != nil {
c.JSON(http.StatusBadRequest, gin.H{"error": err.Error()})
return
}
ch := make(chan any)
go func() {
defer close(ch)
2023-07-18 18:51:30 -07:00
fn := func(r api.ProgressResponse) {
ch <- r
}
2023-07-18 18:51:30 -07:00
if err := PullModel(req.Name, req.Username, req.Password, fn); err != nil {
2023-07-20 12:12:08 -07:00
ch <- gin.H{"error": err.Error()}
}
}()
streamResponse(c, ch)
}
func push(c *gin.Context) {
var req api.PushRequest
if err := c.ShouldBindJSON(&req); err != nil {
c.JSON(http.StatusBadRequest, gin.H{"error": err.Error()})
2023-07-11 11:54:22 -07:00
return
}
2023-07-06 10:40:11 -07:00
ch := make(chan any)
go func() {
defer close(ch)
2023-07-18 18:51:30 -07:00
fn := func(r api.ProgressResponse) {
ch <- r
}
2023-07-18 18:51:30 -07:00
if err := PushModel(req.Name, req.Username, req.Password, fn); err != nil {
2023-07-20 12:12:08 -07:00
ch <- gin.H{"error": err.Error()}
}
}()
streamResponse(c, ch)
}
func create(c *gin.Context) {
var req api.CreateRequest
if err := c.ShouldBindJSON(&req); err != nil {
c.JSON(http.StatusBadRequest, gin.H{"message": err.Error()})
2023-07-12 19:07:15 -07:00
return
}
2023-07-11 11:54:22 -07:00
ch := make(chan any)
go func() {
defer close(ch)
fn := func(status string) {
ch <- api.CreateProgress{
Status: status,
}
}
if err := CreateModel(req.Name, req.Path, fn); err != nil {
2023-07-20 12:12:08 -07:00
ch <- gin.H{"error": err.Error()}
}
}()
2023-07-07 15:29:17 -07:00
streamResponse(c, ch)
2023-07-05 15:37:33 -04:00
}
2023-07-18 09:09:45 -07:00
func list(c *gin.Context) {
var models []api.ListResponseModel
fp, err := GetManifestPath()
if err != nil {
c.JSON(http.StatusInternalServerError, gin.H{"error": err.Error()})
return
}
err = filepath.Walk(fp, func(path string, info os.FileInfo, err error) error {
if err != nil {
return err
}
if !info.IsDir() {
fi, err := os.Stat(path)
if err != nil {
log.Printf("skipping file: %s", fp)
return nil
2023-07-18 09:09:45 -07:00
}
path := path[len(fp)+1:]
slashIndex := strings.LastIndex(path, "/")
if slashIndex == -1 {
return nil
}
tag := path[:slashIndex] + ":" + path[slashIndex+1:]
mp := ParseModelPath(tag)
manifest, err := GetManifest(mp)
if err != nil {
log.Printf("skipping file: %s", fp)
return nil
2023-07-18 09:09:45 -07:00
}
model := api.ListResponseModel{
Name: mp.GetShortTagname(),
Size: manifest.GetTotalSize(),
ModifiedAt: fi.ModTime(),
}
models = append(models, model)
}
return nil
})
if err != nil {
c.JSON(http.StatusInternalServerError, gin.H{"error": err.Error()})
return
}
c.JSON(http.StatusOK, api.ListResponse{models})
}
2023-07-05 15:37:33 -04:00
func Serve(ln net.Listener) error {
r := gin.Default()
2023-07-07 23:46:15 -04:00
r.GET("/", func(c *gin.Context) {
c.String(http.StatusOK, "Ollama is running")
})
2023-07-12 17:19:03 -07:00
r.POST("/api/pull", pull)
2023-07-05 15:37:33 -04:00
r.POST("/api/generate", generate)
r.POST("/api/create", create)
r.POST("/api/push", push)
2023-07-18 09:09:45 -07:00
r.GET("/api/tags", list)
log.Printf("Listening on %s", ln.Addr())
s := &http.Server{
Handler: r,
}
return s.Serve(ln)
}
2023-07-06 10:40:11 -07:00
func streamResponse(c *gin.Context, ch chan any) {
2023-07-11 11:54:22 -07:00
c.Stream(func(w io.Writer) bool {
val, ok := <-ch
if !ok {
return false
}
bts, err := json.Marshal(val)
if err != nil {
return false
}
bts = append(bts, '\n')
if _, err := w.Write(bts); err != nil {
return false
}
return true
})
}