| 123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214215216217218219220221222223224225226227228229230231232233234235236237238239240241242243244245246247248249250251252253254255256257258259260261262263264265266267268269270271272273274275276277278279280281282283284285286287288289290291292293294295296297298299300301302303304305306307308309310311312313314315316317318319320321322323324325326327328329330331332333334335336337338339340341342343344345346347348349350351352353354355356357358359360361362363364365366367368369370371372373374375376377378379380381382383 |
- package vertex
- import (
- "encoding/json"
- "errors"
- "fmt"
- "io"
- "net/http"
- "strings"
- "github.com/QuantumNous/new-api/common"
- "github.com/QuantumNous/new-api/dto"
- "github.com/QuantumNous/new-api/relay/channel"
- "github.com/QuantumNous/new-api/relay/channel/claude"
- "github.com/QuantumNous/new-api/relay/channel/gemini"
- "github.com/QuantumNous/new-api/relay/channel/openai"
- relaycommon "github.com/QuantumNous/new-api/relay/common"
- "github.com/QuantumNous/new-api/relay/constant"
- "github.com/QuantumNous/new-api/setting/model_setting"
- "github.com/QuantumNous/new-api/types"
- "github.com/gin-gonic/gin"
- )
- const (
- RequestModeClaude = 1
- RequestModeGemini = 2
- RequestModeLlama = 3
- )
- var claudeModelMap = map[string]string{
- "claude-3-sonnet-20240229": "claude-3-sonnet@20240229",
- "claude-3-opus-20240229": "claude-3-opus@20240229",
- "claude-3-haiku-20240307": "claude-3-haiku@20240307",
- "claude-3-5-sonnet-20240620": "claude-3-5-sonnet@20240620",
- "claude-3-5-sonnet-20241022": "claude-3-5-sonnet-v2@20241022",
- "claude-3-7-sonnet-20250219": "claude-3-7-sonnet@20250219",
- "claude-sonnet-4-20250514": "claude-sonnet-4@20250514",
- "claude-opus-4-20250514": "claude-opus-4@20250514",
- "claude-opus-4-1-20250805": "claude-opus-4-1@20250805",
- "claude-sonnet-4-5-20250929": "claude-sonnet-4-5@20250929",
- "claude-opus-4-5-20251101": "claude-opus-4-5@20251101",
- }
- const anthropicVersion = "vertex-2023-10-16"
- type Adaptor struct {
- RequestMode int
- AccountCredentials Credentials
- }
- func (a *Adaptor) ConvertGeminiRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeminiChatRequest) (any, error) {
- geminiAdaptor := gemini.Adaptor{}
- return geminiAdaptor.ConvertGeminiRequest(c, info, request)
- }
- func (a *Adaptor) ConvertClaudeRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.ClaudeRequest) (any, error) {
- if v, ok := claudeModelMap[info.UpstreamModelName]; ok {
- c.Set("request_model", v)
- } else {
- c.Set("request_model", request.Model)
- }
- vertexClaudeReq := copyRequest(request, anthropicVersion)
- return vertexClaudeReq, nil
- }
- func (a *Adaptor) ConvertAudioRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.AudioRequest) (io.Reader, error) {
- //TODO implement me
- return nil, errors.New("not implemented")
- }
- func (a *Adaptor) ConvertImageRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.ImageRequest) (any, error) {
- geminiAdaptor := gemini.Adaptor{}
- return geminiAdaptor.ConvertImageRequest(c, info, request)
- }
- func (a *Adaptor) Init(info *relaycommon.RelayInfo) {
- if strings.HasPrefix(info.UpstreamModelName, "claude") {
- a.RequestMode = RequestModeClaude
- } else if strings.Contains(info.UpstreamModelName, "llama") ||
- // open source models
- strings.Contains(info.UpstreamModelName, "-maas") {
- a.RequestMode = RequestModeLlama
- } else {
- a.RequestMode = RequestModeGemini
- }
- }
- func (a *Adaptor) getRequestUrl(info *relaycommon.RelayInfo, modelName, suffix string) (string, error) {
- region := GetModelRegion(info.ApiVersion, info.OriginModelName)
- if info.ChannelOtherSettings.VertexKeyType != dto.VertexKeyTypeAPIKey {
- adc := &Credentials{}
- if err := common.Unmarshal([]byte(info.ApiKey), adc); err != nil {
- return "", fmt.Errorf("failed to decode credentials file: %w", err)
- }
- a.AccountCredentials = *adc
- if a.RequestMode == RequestModeGemini {
- if region == "global" {
- return fmt.Sprintf(
- "https://aiplatform.googleapis.com/v1/projects/%s/locations/global/publishers/google/models/%s:%s",
- adc.ProjectID,
- modelName,
- suffix,
- ), nil
- } else {
- return fmt.Sprintf(
- "https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/google/models/%s:%s",
- region,
- adc.ProjectID,
- region,
- modelName,
- suffix,
- ), nil
- }
- } else if a.RequestMode == RequestModeClaude {
- if region == "global" {
- return fmt.Sprintf(
- "https://aiplatform.googleapis.com/v1/projects/%s/locations/global/publishers/anthropic/models/%s:%s",
- adc.ProjectID,
- modelName,
- suffix,
- ), nil
- } else {
- return fmt.Sprintf(
- "https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/anthropic/models/%s:%s",
- region,
- adc.ProjectID,
- region,
- modelName,
- suffix,
- ), nil
- }
- } else if a.RequestMode == RequestModeLlama {
- return fmt.Sprintf(
- "https://%s-aiplatform.googleapis.com/v1beta1/projects/%s/locations/%s/endpoints/openapi/chat/completions",
- region,
- adc.ProjectID,
- region,
- ), nil
- }
- } else {
- var keyPrefix string
- if strings.HasSuffix(suffix, "?alt=sse") {
- keyPrefix = "&"
- } else {
- keyPrefix = "?"
- }
- if region == "global" {
- return fmt.Sprintf(
- "https://aiplatform.googleapis.com/v1/publishers/google/models/%s:%s%skey=%s",
- modelName,
- suffix,
- keyPrefix,
- info.ApiKey,
- ), nil
- } else {
- return fmt.Sprintf(
- "https://%s-aiplatform.googleapis.com/v1/publishers/google/models/%s:%s%skey=%s",
- region,
- modelName,
- suffix,
- keyPrefix,
- info.ApiKey,
- ), nil
- }
- }
- return "", errors.New("unsupported request mode")
- }
- func (a *Adaptor) GetRequestURL(info *relaycommon.RelayInfo) (string, error) {
- suffix := ""
- if a.RequestMode == RequestModeGemini {
- if model_setting.GetGeminiSettings().ThinkingAdapterEnabled &&
- !model_setting.ShouldPreserveThinkingSuffix(info.OriginModelName) {
- // 新增逻辑:处理 -thinking-<budget> 格式
- if strings.Contains(info.UpstreamModelName, "-thinking-") {
- parts := strings.Split(info.UpstreamModelName, "-thinking-")
- info.UpstreamModelName = parts[0]
- } else if strings.HasSuffix(info.UpstreamModelName, "-thinking") { // 旧的适配
- info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-thinking")
- } else if strings.HasSuffix(info.UpstreamModelName, "-nothinking") {
- info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-nothinking")
- }
- }
- if info.IsStream {
- suffix = "streamGenerateContent?alt=sse"
- } else {
- suffix = "generateContent"
- }
- if strings.HasPrefix(info.UpstreamModelName, "imagen") {
- suffix = "predict"
- }
- return a.getRequestUrl(info, info.UpstreamModelName, suffix)
- } else if a.RequestMode == RequestModeClaude {
- if info.IsStream {
- suffix = "streamRawPredict?alt=sse"
- } else {
- suffix = "rawPredict"
- }
- model := info.UpstreamModelName
- if v, ok := claudeModelMap[info.UpstreamModelName]; ok {
- model = v
- }
- return a.getRequestUrl(info, model, suffix)
- } else if a.RequestMode == RequestModeLlama {
- return a.getRequestUrl(info, "", "")
- }
- return "", errors.New("unsupported request mode")
- }
- func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Header, info *relaycommon.RelayInfo) error {
- channel.SetupApiRequestHeader(info, c, req)
- if info.ChannelOtherSettings.VertexKeyType != dto.VertexKeyTypeAPIKey {
- accessToken, err := getAccessToken(a, info)
- if err != nil {
- return err
- }
- req.Set("Authorization", "Bearer "+accessToken)
- }
- if a.AccountCredentials.ProjectID != "" {
- req.Set("x-goog-user-project", a.AccountCredentials.ProjectID)
- }
- if strings.Contains(info.UpstreamModelName, "claude") {
- claude.CommonClaudeHeadersOperation(c, req, info)
- }
- return nil
- }
- func (a *Adaptor) ConvertOpenAIRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeneralOpenAIRequest) (any, error) {
- if request == nil {
- return nil, errors.New("request is nil")
- }
- if a.RequestMode == RequestModeGemini && strings.HasPrefix(info.UpstreamModelName, "imagen") {
- prompt := ""
- for _, m := range request.Messages {
- if m.Role == "user" {
- prompt = m.StringContent()
- if prompt != "" {
- break
- }
- }
- }
- if prompt == "" {
- if p, ok := request.Prompt.(string); ok {
- prompt = p
- }
- }
- if prompt == "" {
- return nil, errors.New("prompt is required for image generation")
- }
- imgReq := dto.ImageRequest{
- Model: request.Model,
- Prompt: prompt,
- N: 1,
- Size: "1024x1024",
- }
- if request.N > 0 {
- imgReq.N = uint(request.N)
- }
- if request.Size != "" {
- imgReq.Size = request.Size
- }
- if len(request.ExtraBody) > 0 {
- var extra map[string]any
- if err := json.Unmarshal(request.ExtraBody, &extra); err == nil {
- if n, ok := extra["n"].(float64); ok && n > 0 {
- imgReq.N = uint(n)
- }
- if size, ok := extra["size"].(string); ok {
- imgReq.Size = size
- }
- // accept aspectRatio in extra body (top-level or under parameters)
- if ar, ok := extra["aspectRatio"].(string); ok && ar != "" {
- imgReq.Size = ar
- }
- if params, ok := extra["parameters"].(map[string]any); ok {
- if ar, ok := params["aspectRatio"].(string); ok && ar != "" {
- imgReq.Size = ar
- }
- }
- }
- }
- c.Set("request_model", request.Model)
- return a.ConvertImageRequest(c, info, imgReq)
- }
- if a.RequestMode == RequestModeClaude {
- claudeReq, err := claude.RequestOpenAI2ClaudeMessage(c, *request)
- if err != nil {
- return nil, err
- }
- vertexClaudeReq := copyRequest(claudeReq, anthropicVersion)
- c.Set("request_model", claudeReq.Model)
- info.UpstreamModelName = claudeReq.Model
- return vertexClaudeReq, nil
- } else if a.RequestMode == RequestModeGemini {
- geminiRequest, err := gemini.CovertOpenAI2Gemini(c, *request, info)
- if err != nil {
- return nil, err
- }
- c.Set("request_model", request.Model)
- return geminiRequest, nil
- } else if a.RequestMode == RequestModeLlama {
- return request, nil
- }
- return nil, errors.New("unsupported request mode")
- }
- func (a *Adaptor) ConvertRerankRequest(c *gin.Context, relayMode int, request dto.RerankRequest) (any, error) {
- return nil, nil
- }
- func (a *Adaptor) ConvertEmbeddingRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.EmbeddingRequest) (any, error) {
- //TODO implement me
- return nil, errors.New("not implemented")
- }
- func (a *Adaptor) ConvertOpenAIResponsesRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.OpenAIResponsesRequest) (any, error) {
- // TODO implement me
- return nil, errors.New("not implemented")
- }
- func (a *Adaptor) DoRequest(c *gin.Context, info *relaycommon.RelayInfo, requestBody io.Reader) (any, error) {
- return channel.DoApiRequest(a, c, info, requestBody)
- }
- func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, info *relaycommon.RelayInfo) (usage any, err *types.NewAPIError) {
- if info.IsStream {
- switch a.RequestMode {
- case RequestModeClaude:
- return claude.ClaudeStreamHandler(c, resp, info, claude.RequestModeMessage)
- case RequestModeGemini:
- if info.RelayMode == constant.RelayModeGemini {
- return gemini.GeminiTextGenerationStreamHandler(c, info, resp)
- } else {
- return gemini.GeminiChatStreamHandler(c, info, resp)
- }
- case RequestModeLlama:
- return openai.OaiStreamHandler(c, info, resp)
- }
- } else {
- switch a.RequestMode {
- case RequestModeClaude:
- return claude.ClaudeHandler(c, resp, info, claude.RequestModeMessage)
- case RequestModeGemini:
- if info.RelayMode == constant.RelayModeGemini {
- return gemini.GeminiTextGenerationHandler(c, info, resp)
- } else {
- if strings.HasPrefix(info.UpstreamModelName, "imagen") {
- return gemini.GeminiImageHandler(c, info, resp)
- }
- return gemini.GeminiChatHandler(c, info, resp)
- }
- case RequestModeLlama:
- return openai.OpenaiHandler(c, info, resp)
- }
- }
- return
- }
- func (a *Adaptor) GetModelList() []string {
- var modelList []string
- for i, s := range ModelList {
- modelList = append(modelList, s)
- ModelList[i] = s
- }
- for i, s := range claude.ModelList {
- modelList = append(modelList, s)
- claude.ModelList[i] = s
- }
- for i, s := range gemini.ModelList {
- modelList = append(modelList, s)
- gemini.ModelList[i] = s
- }
- return modelList
- }
- func (a *Adaptor) GetChannelName() string {
- return ChannelName
- }
|