adaptor.go 12 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214215216217218219220221222223224225226227228229230231232233234235236237238239240241242243244245246247248249250251252253254255256257258259260261262263264265266267268269270271272273274275276277278279280281282283284285286287288289290291292293294295296297298299300301302303304305306307308309310311312313314315316317318319320321322323324325326327328329330331332333334335336337338339340341342343344345346347348349350351352353354355356357358359360361362363364365366367368369370371372373374375376377378379380381382383384385386
  1. package vertex
  2. import (
  3. "encoding/json"
  4. "errors"
  5. "fmt"
  6. "io"
  7. "net/http"
  8. "strings"
  9. "github.com/QuantumNous/new-api/common"
  10. "github.com/QuantumNous/new-api/dto"
  11. "github.com/QuantumNous/new-api/relay/channel"
  12. "github.com/QuantumNous/new-api/relay/channel/claude"
  13. "github.com/QuantumNous/new-api/relay/channel/gemini"
  14. "github.com/QuantumNous/new-api/relay/channel/openai"
  15. relaycommon "github.com/QuantumNous/new-api/relay/common"
  16. "github.com/QuantumNous/new-api/relay/constant"
  17. "github.com/QuantumNous/new-api/setting/model_setting"
  18. "github.com/QuantumNous/new-api/setting/reasoning"
  19. "github.com/QuantumNous/new-api/types"
  20. "github.com/gin-gonic/gin"
  21. )
  22. const (
  23. RequestModeClaude = 1
  24. RequestModeGemini = 2
  25. RequestModeLlama = 3
  26. )
  27. var claudeModelMap = map[string]string{
  28. "claude-3-sonnet-20240229": "claude-3-sonnet@20240229",
  29. "claude-3-opus-20240229": "claude-3-opus@20240229",
  30. "claude-3-haiku-20240307": "claude-3-haiku@20240307",
  31. "claude-3-5-sonnet-20240620": "claude-3-5-sonnet@20240620",
  32. "claude-3-5-sonnet-20241022": "claude-3-5-sonnet-v2@20241022",
  33. "claude-3-7-sonnet-20250219": "claude-3-7-sonnet@20250219",
  34. "claude-sonnet-4-20250514": "claude-sonnet-4@20250514",
  35. "claude-opus-4-20250514": "claude-opus-4@20250514",
  36. "claude-opus-4-1-20250805": "claude-opus-4-1@20250805",
  37. "claude-sonnet-4-5-20250929": "claude-sonnet-4-5@20250929",
  38. "claude-opus-4-5-20251101": "claude-opus-4-5@20251101",
  39. }
  40. const anthropicVersion = "vertex-2023-10-16"
  41. type Adaptor struct {
  42. RequestMode int
  43. AccountCredentials Credentials
  44. }
  45. func (a *Adaptor) ConvertGeminiRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeminiChatRequest) (any, error) {
  46. geminiAdaptor := gemini.Adaptor{}
  47. return geminiAdaptor.ConvertGeminiRequest(c, info, request)
  48. }
  49. func (a *Adaptor) ConvertClaudeRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.ClaudeRequest) (any, error) {
  50. if v, ok := claudeModelMap[info.UpstreamModelName]; ok {
  51. c.Set("request_model", v)
  52. } else {
  53. c.Set("request_model", request.Model)
  54. }
  55. vertexClaudeReq := copyRequest(request, anthropicVersion)
  56. return vertexClaudeReq, nil
  57. }
  58. func (a *Adaptor) ConvertAudioRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.AudioRequest) (io.Reader, error) {
  59. //TODO implement me
  60. return nil, errors.New("not implemented")
  61. }
  62. func (a *Adaptor) ConvertImageRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.ImageRequest) (any, error) {
  63. geminiAdaptor := gemini.Adaptor{}
  64. return geminiAdaptor.ConvertImageRequest(c, info, request)
  65. }
  66. func (a *Adaptor) Init(info *relaycommon.RelayInfo) {
  67. if strings.HasPrefix(info.UpstreamModelName, "claude") {
  68. a.RequestMode = RequestModeClaude
  69. } else if strings.Contains(info.UpstreamModelName, "llama") ||
  70. // open source models
  71. strings.Contains(info.UpstreamModelName, "-maas") {
  72. a.RequestMode = RequestModeLlama
  73. } else {
  74. a.RequestMode = RequestModeGemini
  75. }
  76. }
  77. func (a *Adaptor) getRequestUrl(info *relaycommon.RelayInfo, modelName, suffix string) (string, error) {
  78. region := GetModelRegion(info.ApiVersion, info.OriginModelName)
  79. if info.ChannelOtherSettings.VertexKeyType != dto.VertexKeyTypeAPIKey {
  80. adc := &Credentials{}
  81. if err := common.Unmarshal([]byte(info.ApiKey), adc); err != nil {
  82. return "", fmt.Errorf("failed to decode credentials file: %w", err)
  83. }
  84. a.AccountCredentials = *adc
  85. if a.RequestMode == RequestModeGemini {
  86. if region == "global" {
  87. return fmt.Sprintf(
  88. "https://aiplatform.googleapis.com/v1/projects/%s/locations/global/publishers/google/models/%s:%s",
  89. adc.ProjectID,
  90. modelName,
  91. suffix,
  92. ), nil
  93. } else {
  94. return fmt.Sprintf(
  95. "https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/google/models/%s:%s",
  96. region,
  97. adc.ProjectID,
  98. region,
  99. modelName,
  100. suffix,
  101. ), nil
  102. }
  103. } else if a.RequestMode == RequestModeClaude {
  104. if region == "global" {
  105. return fmt.Sprintf(
  106. "https://aiplatform.googleapis.com/v1/projects/%s/locations/global/publishers/anthropic/models/%s:%s",
  107. adc.ProjectID,
  108. modelName,
  109. suffix,
  110. ), nil
  111. } else {
  112. return fmt.Sprintf(
  113. "https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/anthropic/models/%s:%s",
  114. region,
  115. adc.ProjectID,
  116. region,
  117. modelName,
  118. suffix,
  119. ), nil
  120. }
  121. } else if a.RequestMode == RequestModeLlama {
  122. return fmt.Sprintf(
  123. "https://%s-aiplatform.googleapis.com/v1beta1/projects/%s/locations/%s/endpoints/openapi/chat/completions",
  124. region,
  125. adc.ProjectID,
  126. region,
  127. ), nil
  128. }
  129. } else {
  130. var keyPrefix string
  131. if strings.HasSuffix(suffix, "?alt=sse") {
  132. keyPrefix = "&"
  133. } else {
  134. keyPrefix = "?"
  135. }
  136. if region == "global" {
  137. return fmt.Sprintf(
  138. "https://aiplatform.googleapis.com/v1/publishers/google/models/%s:%s%skey=%s",
  139. modelName,
  140. suffix,
  141. keyPrefix,
  142. info.ApiKey,
  143. ), nil
  144. } else {
  145. return fmt.Sprintf(
  146. "https://%s-aiplatform.googleapis.com/v1/publishers/google/models/%s:%s%skey=%s",
  147. region,
  148. modelName,
  149. suffix,
  150. keyPrefix,
  151. info.ApiKey,
  152. ), nil
  153. }
  154. }
  155. return "", errors.New("unsupported request mode")
  156. }
  157. func (a *Adaptor) GetRequestURL(info *relaycommon.RelayInfo) (string, error) {
  158. suffix := ""
  159. if a.RequestMode == RequestModeGemini {
  160. if model_setting.GetGeminiSettings().ThinkingAdapterEnabled &&
  161. !model_setting.ShouldPreserveThinkingSuffix(info.OriginModelName) {
  162. // 新增逻辑:处理 -thinking-<budget> 格式
  163. if strings.Contains(info.UpstreamModelName, "-thinking-") {
  164. parts := strings.Split(info.UpstreamModelName, "-thinking-")
  165. info.UpstreamModelName = parts[0]
  166. } else if strings.HasSuffix(info.UpstreamModelName, "-thinking") { // 旧的适配
  167. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-thinking")
  168. } else if strings.HasSuffix(info.UpstreamModelName, "-nothinking") {
  169. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-nothinking")
  170. } else if baseModel, level, ok := reasoning.TrimEffortSuffix(info.UpstreamModelName); ok && level != "" {
  171. info.UpstreamModelName = baseModel
  172. }
  173. }
  174. if info.IsStream {
  175. suffix = "streamGenerateContent?alt=sse"
  176. } else {
  177. suffix = "generateContent"
  178. }
  179. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  180. suffix = "predict"
  181. }
  182. return a.getRequestUrl(info, info.UpstreamModelName, suffix)
  183. } else if a.RequestMode == RequestModeClaude {
  184. if info.IsStream {
  185. suffix = "streamRawPredict?alt=sse"
  186. } else {
  187. suffix = "rawPredict"
  188. }
  189. model := info.UpstreamModelName
  190. if v, ok := claudeModelMap[info.UpstreamModelName]; ok {
  191. model = v
  192. }
  193. return a.getRequestUrl(info, model, suffix)
  194. } else if a.RequestMode == RequestModeLlama {
  195. return a.getRequestUrl(info, "", "")
  196. }
  197. return "", errors.New("unsupported request mode")
  198. }
  199. func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Header, info *relaycommon.RelayInfo) error {
  200. channel.SetupApiRequestHeader(info, c, req)
  201. if info.ChannelOtherSettings.VertexKeyType != dto.VertexKeyTypeAPIKey {
  202. accessToken, err := getAccessToken(a, info)
  203. if err != nil {
  204. return err
  205. }
  206. req.Set("Authorization", "Bearer "+accessToken)
  207. }
  208. if a.AccountCredentials.ProjectID != "" {
  209. req.Set("x-goog-user-project", a.AccountCredentials.ProjectID)
  210. }
  211. if strings.Contains(info.UpstreamModelName, "claude") {
  212. claude.CommonClaudeHeadersOperation(c, req, info)
  213. }
  214. return nil
  215. }
  216. func (a *Adaptor) ConvertOpenAIRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeneralOpenAIRequest) (any, error) {
  217. if request == nil {
  218. return nil, errors.New("request is nil")
  219. }
  220. if a.RequestMode == RequestModeGemini && strings.HasPrefix(info.UpstreamModelName, "imagen") {
  221. prompt := ""
  222. for _, m := range request.Messages {
  223. if m.Role == "user" {
  224. prompt = m.StringContent()
  225. if prompt != "" {
  226. break
  227. }
  228. }
  229. }
  230. if prompt == "" {
  231. if p, ok := request.Prompt.(string); ok {
  232. prompt = p
  233. }
  234. }
  235. if prompt == "" {
  236. return nil, errors.New("prompt is required for image generation")
  237. }
  238. imgReq := dto.ImageRequest{
  239. Model: request.Model,
  240. Prompt: prompt,
  241. N: 1,
  242. Size: "1024x1024",
  243. }
  244. if request.N > 0 {
  245. imgReq.N = uint(request.N)
  246. }
  247. if request.Size != "" {
  248. imgReq.Size = request.Size
  249. }
  250. if len(request.ExtraBody) > 0 {
  251. var extra map[string]any
  252. if err := json.Unmarshal(request.ExtraBody, &extra); err == nil {
  253. if n, ok := extra["n"].(float64); ok && n > 0 {
  254. imgReq.N = uint(n)
  255. }
  256. if size, ok := extra["size"].(string); ok {
  257. imgReq.Size = size
  258. }
  259. // accept aspectRatio in extra body (top-level or under parameters)
  260. if ar, ok := extra["aspectRatio"].(string); ok && ar != "" {
  261. imgReq.Size = ar
  262. }
  263. if params, ok := extra["parameters"].(map[string]any); ok {
  264. if ar, ok := params["aspectRatio"].(string); ok && ar != "" {
  265. imgReq.Size = ar
  266. }
  267. }
  268. }
  269. }
  270. c.Set("request_model", request.Model)
  271. return a.ConvertImageRequest(c, info, imgReq)
  272. }
  273. if a.RequestMode == RequestModeClaude {
  274. claudeReq, err := claude.RequestOpenAI2ClaudeMessage(c, *request)
  275. if err != nil {
  276. return nil, err
  277. }
  278. vertexClaudeReq := copyRequest(claudeReq, anthropicVersion)
  279. c.Set("request_model", claudeReq.Model)
  280. info.UpstreamModelName = claudeReq.Model
  281. return vertexClaudeReq, nil
  282. } else if a.RequestMode == RequestModeGemini {
  283. geminiRequest, err := gemini.CovertOpenAI2Gemini(c, *request, info)
  284. if err != nil {
  285. return nil, err
  286. }
  287. c.Set("request_model", request.Model)
  288. return geminiRequest, nil
  289. } else if a.RequestMode == RequestModeLlama {
  290. return request, nil
  291. }
  292. return nil, errors.New("unsupported request mode")
  293. }
  294. func (a *Adaptor) ConvertRerankRequest(c *gin.Context, relayMode int, request dto.RerankRequest) (any, error) {
  295. return nil, nil
  296. }
  297. func (a *Adaptor) ConvertEmbeddingRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.EmbeddingRequest) (any, error) {
  298. //TODO implement me
  299. return nil, errors.New("not implemented")
  300. }
  301. func (a *Adaptor) ConvertOpenAIResponsesRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.OpenAIResponsesRequest) (any, error) {
  302. // TODO implement me
  303. return nil, errors.New("not implemented")
  304. }
  305. func (a *Adaptor) DoRequest(c *gin.Context, info *relaycommon.RelayInfo, requestBody io.Reader) (any, error) {
  306. return channel.DoApiRequest(a, c, info, requestBody)
  307. }
  308. func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, info *relaycommon.RelayInfo) (usage any, err *types.NewAPIError) {
  309. if info.IsStream {
  310. switch a.RequestMode {
  311. case RequestModeClaude:
  312. return claude.ClaudeStreamHandler(c, resp, info, claude.RequestModeMessage)
  313. case RequestModeGemini:
  314. if info.RelayMode == constant.RelayModeGemini {
  315. return gemini.GeminiTextGenerationStreamHandler(c, info, resp)
  316. } else {
  317. return gemini.GeminiChatStreamHandler(c, info, resp)
  318. }
  319. case RequestModeLlama:
  320. return openai.OaiStreamHandler(c, info, resp)
  321. }
  322. } else {
  323. switch a.RequestMode {
  324. case RequestModeClaude:
  325. return claude.ClaudeHandler(c, resp, info, claude.RequestModeMessage)
  326. case RequestModeGemini:
  327. if info.RelayMode == constant.RelayModeGemini {
  328. return gemini.GeminiTextGenerationHandler(c, info, resp)
  329. } else {
  330. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  331. return gemini.GeminiImageHandler(c, info, resp)
  332. }
  333. return gemini.GeminiChatHandler(c, info, resp)
  334. }
  335. case RequestModeLlama:
  336. return openai.OpenaiHandler(c, info, resp)
  337. }
  338. }
  339. return
  340. }
  341. func (a *Adaptor) GetModelList() []string {
  342. var modelList []string
  343. for i, s := range ModelList {
  344. modelList = append(modelList, s)
  345. ModelList[i] = s
  346. }
  347. for i, s := range claude.ModelList {
  348. modelList = append(modelList, s)
  349. claude.ModelList[i] = s
  350. }
  351. for i, s := range gemini.ModelList {
  352. modelList = append(modelList, s)
  353. gemini.ModelList[i] = s
  354. }
  355. return modelList
  356. }
  357. func (a *Adaptor) GetChannelName() string {
  358. return ChannelName
  359. }