adaptor.go 11 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214215216217218219220221222223224225226227228229230231232233234235236237238239240241242243244245246247248249250251252253254255256257258259260261262263264265266267268269270271272273274275276277278279280281282283284285286287288289290291292293294295296297298299300301302303304305306307308309310311312313314315316317318319320321322323324325326327328329330331332333334335336337338339340341342343344345346347348349350351352353354355356357358359360361362363364365366367368369370371372373374375376
  1. package vertex
  2. import (
  3. "encoding/json"
  4. "errors"
  5. "fmt"
  6. "io"
  7. "net/http"
  8. "strings"
  9. "github.com/QuantumNous/new-api/common"
  10. "github.com/QuantumNous/new-api/dto"
  11. "github.com/QuantumNous/new-api/relay/channel"
  12. "github.com/QuantumNous/new-api/relay/channel/claude"
  13. "github.com/QuantumNous/new-api/relay/channel/gemini"
  14. "github.com/QuantumNous/new-api/relay/channel/openai"
  15. relaycommon "github.com/QuantumNous/new-api/relay/common"
  16. "github.com/QuantumNous/new-api/relay/constant"
  17. "github.com/QuantumNous/new-api/setting/model_setting"
  18. "github.com/QuantumNous/new-api/types"
  19. "github.com/gin-gonic/gin"
  20. )
  21. const (
  22. RequestModeClaude = 1
  23. RequestModeGemini = 2
  24. RequestModeLlama = 3
  25. )
  26. var claudeModelMap = map[string]string{
  27. "claude-3-sonnet-20240229": "claude-3-sonnet@20240229",
  28. "claude-3-opus-20240229": "claude-3-opus@20240229",
  29. "claude-3-haiku-20240307": "claude-3-haiku@20240307",
  30. "claude-3-5-sonnet-20240620": "claude-3-5-sonnet@20240620",
  31. "claude-3-5-sonnet-20241022": "claude-3-5-sonnet-v2@20241022",
  32. "claude-3-7-sonnet-20250219": "claude-3-7-sonnet@20250219",
  33. "claude-sonnet-4-20250514": "claude-sonnet-4@20250514",
  34. "claude-opus-4-20250514": "claude-opus-4@20250514",
  35. "claude-opus-4-1-20250805": "claude-opus-4-1@20250805",
  36. "claude-sonnet-4-5-20250929": "claude-sonnet-4-5@20250929",
  37. }
  38. const anthropicVersion = "vertex-2023-10-16"
  39. type Adaptor struct {
  40. RequestMode int
  41. AccountCredentials Credentials
  42. }
  43. func (a *Adaptor) ConvertGeminiRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeminiChatRequest) (any, error) {
  44. geminiAdaptor := gemini.Adaptor{}
  45. return geminiAdaptor.ConvertGeminiRequest(c, info, request)
  46. }
  47. func (a *Adaptor) ConvertClaudeRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.ClaudeRequest) (any, error) {
  48. if v, ok := claudeModelMap[info.UpstreamModelName]; ok {
  49. c.Set("request_model", v)
  50. } else {
  51. c.Set("request_model", request.Model)
  52. }
  53. vertexClaudeReq := copyRequest(request, anthropicVersion)
  54. return vertexClaudeReq, nil
  55. }
  56. func (a *Adaptor) ConvertAudioRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.AudioRequest) (io.Reader, error) {
  57. //TODO implement me
  58. return nil, errors.New("not implemented")
  59. }
  60. func (a *Adaptor) ConvertImageRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.ImageRequest) (any, error) {
  61. geminiAdaptor := gemini.Adaptor{}
  62. return geminiAdaptor.ConvertImageRequest(c, info, request)
  63. }
  64. func (a *Adaptor) Init(info *relaycommon.RelayInfo) {
  65. if strings.HasPrefix(info.UpstreamModelName, "claude") {
  66. a.RequestMode = RequestModeClaude
  67. } else if strings.Contains(info.UpstreamModelName, "llama") {
  68. a.RequestMode = RequestModeLlama
  69. } else {
  70. a.RequestMode = RequestModeGemini
  71. }
  72. }
  73. func (a *Adaptor) getRequestUrl(info *relaycommon.RelayInfo, modelName, suffix string) (string, error) {
  74. region := GetModelRegion(info.ApiVersion, info.OriginModelName)
  75. if info.ChannelOtherSettings.VertexKeyType != dto.VertexKeyTypeAPIKey {
  76. adc := &Credentials{}
  77. if err := common.Unmarshal([]byte(info.ApiKey), adc); err != nil {
  78. return "", fmt.Errorf("failed to decode credentials file: %w", err)
  79. }
  80. a.AccountCredentials = *adc
  81. if a.RequestMode == RequestModeGemini {
  82. if region == "global" {
  83. return fmt.Sprintf(
  84. "https://aiplatform.googleapis.com/v1/projects/%s/locations/global/publishers/google/models/%s:%s",
  85. adc.ProjectID,
  86. modelName,
  87. suffix,
  88. ), nil
  89. } else {
  90. return fmt.Sprintf(
  91. "https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/google/models/%s:%s",
  92. region,
  93. adc.ProjectID,
  94. region,
  95. modelName,
  96. suffix,
  97. ), nil
  98. }
  99. } else if a.RequestMode == RequestModeClaude {
  100. if region == "global" {
  101. return fmt.Sprintf(
  102. "https://aiplatform.googleapis.com/v1/projects/%s/locations/global/publishers/anthropic/models/%s:%s",
  103. adc.ProjectID,
  104. modelName,
  105. suffix,
  106. ), nil
  107. } else {
  108. return fmt.Sprintf(
  109. "https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/anthropic/models/%s:%s",
  110. region,
  111. adc.ProjectID,
  112. region,
  113. modelName,
  114. suffix,
  115. ), nil
  116. }
  117. } else if a.RequestMode == RequestModeLlama {
  118. return fmt.Sprintf(
  119. "https://%s-aiplatform.googleapis.com/v1beta1/projects/%s/locations/%s/endpoints/openapi/chat/completions",
  120. region,
  121. adc.ProjectID,
  122. region,
  123. ), nil
  124. }
  125. } else {
  126. var keyPrefix string
  127. if strings.HasSuffix(suffix, "?alt=sse") {
  128. keyPrefix = "&"
  129. } else {
  130. keyPrefix = "?"
  131. }
  132. if region == "global" {
  133. return fmt.Sprintf(
  134. "https://aiplatform.googleapis.com/v1/publishers/google/models/%s:%s%skey=%s",
  135. modelName,
  136. suffix,
  137. keyPrefix,
  138. info.ApiKey,
  139. ), nil
  140. } else {
  141. return fmt.Sprintf(
  142. "https://%s-aiplatform.googleapis.com/v1/publishers/google/models/%s:%s%skey=%s",
  143. region,
  144. modelName,
  145. suffix,
  146. keyPrefix,
  147. info.ApiKey,
  148. ), nil
  149. }
  150. }
  151. return "", errors.New("unsupported request mode")
  152. }
  153. func (a *Adaptor) GetRequestURL(info *relaycommon.RelayInfo) (string, error) {
  154. suffix := ""
  155. if a.RequestMode == RequestModeGemini {
  156. if model_setting.GetGeminiSettings().ThinkingAdapterEnabled {
  157. // 新增逻辑:处理 -thinking-<budget> 格式
  158. if strings.Contains(info.UpstreamModelName, "-thinking-") {
  159. parts := strings.Split(info.UpstreamModelName, "-thinking-")
  160. info.UpstreamModelName = parts[0]
  161. } else if strings.HasSuffix(info.UpstreamModelName, "-thinking") { // 旧的适配
  162. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-thinking")
  163. } else if strings.HasSuffix(info.UpstreamModelName, "-nothinking") {
  164. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-nothinking")
  165. }
  166. }
  167. if info.IsStream {
  168. suffix = "streamGenerateContent?alt=sse"
  169. } else {
  170. suffix = "generateContent"
  171. }
  172. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  173. suffix = "predict"
  174. }
  175. return a.getRequestUrl(info, info.UpstreamModelName, suffix)
  176. } else if a.RequestMode == RequestModeClaude {
  177. if info.IsStream {
  178. suffix = "streamRawPredict?alt=sse"
  179. } else {
  180. suffix = "rawPredict"
  181. }
  182. model := info.UpstreamModelName
  183. if v, ok := claudeModelMap[info.UpstreamModelName]; ok {
  184. model = v
  185. }
  186. return a.getRequestUrl(info, model, suffix)
  187. } else if a.RequestMode == RequestModeLlama {
  188. return a.getRequestUrl(info, "", "")
  189. }
  190. return "", errors.New("unsupported request mode")
  191. }
  192. func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Header, info *relaycommon.RelayInfo) error {
  193. channel.SetupApiRequestHeader(info, c, req)
  194. if info.ChannelOtherSettings.VertexKeyType != dto.VertexKeyTypeAPIKey {
  195. accessToken, err := getAccessToken(a, info)
  196. if err != nil {
  197. return err
  198. }
  199. req.Set("Authorization", "Bearer "+accessToken)
  200. }
  201. if a.AccountCredentials.ProjectID != "" {
  202. req.Set("x-goog-user-project", a.AccountCredentials.ProjectID)
  203. }
  204. return nil
  205. }
  206. func (a *Adaptor) ConvertOpenAIRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeneralOpenAIRequest) (any, error) {
  207. if request == nil {
  208. return nil, errors.New("request is nil")
  209. }
  210. if a.RequestMode == RequestModeGemini && strings.HasPrefix(info.UpstreamModelName, "imagen") {
  211. prompt := ""
  212. for _, m := range request.Messages {
  213. if m.Role == "user" {
  214. prompt = m.StringContent()
  215. if prompt != "" {
  216. break
  217. }
  218. }
  219. }
  220. if prompt == "" {
  221. if p, ok := request.Prompt.(string); ok {
  222. prompt = p
  223. }
  224. }
  225. if prompt == "" {
  226. return nil, errors.New("prompt is required for image generation")
  227. }
  228. imgReq := dto.ImageRequest{
  229. Model: request.Model,
  230. Prompt: prompt,
  231. N: 1,
  232. Size: "1024x1024",
  233. }
  234. if request.N > 0 {
  235. imgReq.N = uint(request.N)
  236. }
  237. if request.Size != "" {
  238. imgReq.Size = request.Size
  239. }
  240. if len(request.ExtraBody) > 0 {
  241. var extra map[string]any
  242. if err := json.Unmarshal(request.ExtraBody, &extra); err == nil {
  243. if n, ok := extra["n"].(float64); ok && n > 0 {
  244. imgReq.N = uint(n)
  245. }
  246. if size, ok := extra["size"].(string); ok {
  247. imgReq.Size = size
  248. }
  249. // accept aspectRatio in extra body (top-level or under parameters)
  250. if ar, ok := extra["aspectRatio"].(string); ok && ar != "" {
  251. imgReq.Size = ar
  252. }
  253. if params, ok := extra["parameters"].(map[string]any); ok {
  254. if ar, ok := params["aspectRatio"].(string); ok && ar != "" {
  255. imgReq.Size = ar
  256. }
  257. }
  258. }
  259. }
  260. c.Set("request_model", request.Model)
  261. return a.ConvertImageRequest(c, info, imgReq)
  262. }
  263. if a.RequestMode == RequestModeClaude {
  264. claudeReq, err := claude.RequestOpenAI2ClaudeMessage(c, *request)
  265. if err != nil {
  266. return nil, err
  267. }
  268. vertexClaudeReq := copyRequest(claudeReq, anthropicVersion)
  269. c.Set("request_model", claudeReq.Model)
  270. info.UpstreamModelName = claudeReq.Model
  271. return vertexClaudeReq, nil
  272. } else if a.RequestMode == RequestModeGemini {
  273. geminiRequest, err := gemini.CovertGemini2OpenAI(c, *request, info)
  274. if err != nil {
  275. return nil, err
  276. }
  277. c.Set("request_model", request.Model)
  278. return geminiRequest, nil
  279. } else if a.RequestMode == RequestModeLlama {
  280. return request, nil
  281. }
  282. return nil, errors.New("unsupported request mode")
  283. }
  284. func (a *Adaptor) ConvertRerankRequest(c *gin.Context, relayMode int, request dto.RerankRequest) (any, error) {
  285. return nil, nil
  286. }
  287. func (a *Adaptor) ConvertEmbeddingRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.EmbeddingRequest) (any, error) {
  288. //TODO implement me
  289. return nil, errors.New("not implemented")
  290. }
  291. func (a *Adaptor) ConvertOpenAIResponsesRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.OpenAIResponsesRequest) (any, error) {
  292. // TODO implement me
  293. return nil, errors.New("not implemented")
  294. }
  295. func (a *Adaptor) DoRequest(c *gin.Context, info *relaycommon.RelayInfo, requestBody io.Reader) (any, error) {
  296. return channel.DoApiRequest(a, c, info, requestBody)
  297. }
  298. func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, info *relaycommon.RelayInfo) (usage any, err *types.NewAPIError) {
  299. if info.IsStream {
  300. switch a.RequestMode {
  301. case RequestModeClaude:
  302. return claude.ClaudeStreamHandler(c, resp, info, claude.RequestModeMessage)
  303. case RequestModeGemini:
  304. if info.RelayMode == constant.RelayModeGemini {
  305. return gemini.GeminiTextGenerationStreamHandler(c, info, resp)
  306. } else {
  307. return gemini.GeminiChatStreamHandler(c, info, resp)
  308. }
  309. case RequestModeLlama:
  310. return openai.OaiStreamHandler(c, info, resp)
  311. }
  312. } else {
  313. switch a.RequestMode {
  314. case RequestModeClaude:
  315. return claude.ClaudeHandler(c, resp, info, claude.RequestModeMessage)
  316. case RequestModeGemini:
  317. if info.RelayMode == constant.RelayModeGemini {
  318. return gemini.GeminiTextGenerationHandler(c, info, resp)
  319. } else {
  320. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  321. return gemini.GeminiImageHandler(c, info, resp)
  322. }
  323. return gemini.GeminiChatHandler(c, info, resp)
  324. }
  325. case RequestModeLlama:
  326. return openai.OpenaiHandler(c, info, resp)
  327. }
  328. }
  329. return
  330. }
  331. func (a *Adaptor) GetModelList() []string {
  332. var modelList []string
  333. for i, s := range ModelList {
  334. modelList = append(modelList, s)
  335. ModelList[i] = s
  336. }
  337. for i, s := range claude.ModelList {
  338. modelList = append(modelList, s)
  339. claude.ModelList[i] = s
  340. }
  341. for i, s := range gemini.ModelList {
  342. modelList = append(modelList, s)
  343. gemini.ModelList[i] = s
  344. }
  345. return modelList
  346. }
  347. func (a *Adaptor) GetChannelName() string {
  348. return ChannelName
  349. }