adaptor.go 7.6 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214215216217218219220221222223224225226227228229230231232233234235236237238239240241242243244245246247248249250251252253254
  1. package gemini
  2. import (
  3. "errors"
  4. "fmt"
  5. "io"
  6. "net/http"
  7. "one-api/dto"
  8. "one-api/relay/channel"
  9. "one-api/relay/channel/openai"
  10. relaycommon "one-api/relay/common"
  11. "one-api/relay/constant"
  12. "one-api/setting/model_setting"
  13. "one-api/types"
  14. "strings"
  15. "github.com/gin-gonic/gin"
  16. )
  17. type Adaptor struct {
  18. }
  19. func (a *Adaptor) ConvertGeminiRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeminiChatRequest) (any, error) {
  20. if len(request.Contents) > 0 {
  21. for i, content := range request.Contents {
  22. if i == 0 {
  23. if request.Contents[0].Role == "" {
  24. request.Contents[0].Role = "user"
  25. }
  26. }
  27. for _, part := range content.Parts {
  28. if part.FileData != nil {
  29. if part.FileData.MimeType == "" && strings.Contains(part.FileData.FileUri, "www.youtube.com") {
  30. part.FileData.MimeType = "video/webm"
  31. }
  32. }
  33. }
  34. }
  35. }
  36. return request, nil
  37. }
  38. func (a *Adaptor) ConvertClaudeRequest(c *gin.Context, info *relaycommon.RelayInfo, req *dto.ClaudeRequest) (any, error) {
  39. adaptor := openai.Adaptor{}
  40. oaiReq, err := adaptor.ConvertClaudeRequest(c, info, req)
  41. if err != nil {
  42. return nil, err
  43. }
  44. return a.ConvertOpenAIRequest(c, info, oaiReq.(*dto.GeneralOpenAIRequest))
  45. }
  46. func (a *Adaptor) ConvertAudioRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.AudioRequest) (io.Reader, error) {
  47. //TODO implement me
  48. return nil, errors.New("not implemented")
  49. }
  50. func (a *Adaptor) ConvertImageRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.ImageRequest) (any, error) {
  51. if !strings.HasPrefix(info.UpstreamModelName, "imagen") {
  52. return nil, errors.New("not supported model for image generation")
  53. }
  54. // convert size to aspect ratio but allow user to specify aspect ratio
  55. aspectRatio := "1:1" // default aspect ratio
  56. size := strings.TrimSpace(request.Size)
  57. if size != "" {
  58. if strings.Contains(size, ":") {
  59. aspectRatio = size
  60. } else {
  61. switch size {
  62. case "1024x1024":
  63. aspectRatio = "1:1"
  64. case "1024x1792":
  65. aspectRatio = "9:16"
  66. case "1792x1024":
  67. aspectRatio = "16:9"
  68. }
  69. }
  70. }
  71. // build gemini imagen request
  72. geminiRequest := dto.GeminiImageRequest{
  73. Instances: []dto.GeminiImageInstance{
  74. {
  75. Prompt: request.Prompt,
  76. },
  77. },
  78. Parameters: dto.GeminiImageParameters{
  79. SampleCount: int(request.N),
  80. AspectRatio: aspectRatio,
  81. PersonGeneration: "allow_adult", // default allow adult
  82. },
  83. }
  84. return geminiRequest, nil
  85. }
  86. func (a *Adaptor) Init(info *relaycommon.RelayInfo) {
  87. }
  88. func (a *Adaptor) GetRequestURL(info *relaycommon.RelayInfo) (string, error) {
  89. if model_setting.GetGeminiSettings().ThinkingAdapterEnabled {
  90. // 新增逻辑:处理 -thinking-<budget> 格式
  91. if strings.Contains(info.UpstreamModelName, "-thinking-") {
  92. parts := strings.Split(info.UpstreamModelName, "-thinking-")
  93. info.UpstreamModelName = parts[0]
  94. } else if strings.HasSuffix(info.UpstreamModelName, "-thinking") { // 旧的适配
  95. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-thinking")
  96. } else if strings.HasSuffix(info.UpstreamModelName, "-nothinking") {
  97. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-nothinking")
  98. }
  99. }
  100. version := model_setting.GetGeminiVersionSetting(info.UpstreamModelName)
  101. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  102. return fmt.Sprintf("%s/%s/models/%s:predict", info.ChannelBaseUrl, version, info.UpstreamModelName), nil
  103. }
  104. if strings.HasPrefix(info.UpstreamModelName, "text-embedding") ||
  105. strings.HasPrefix(info.UpstreamModelName, "embedding") ||
  106. strings.HasPrefix(info.UpstreamModelName, "gemini-embedding") {
  107. action := "embedContent"
  108. if info.IsGeminiBatchEmbedding {
  109. action = "batchEmbedContents"
  110. }
  111. return fmt.Sprintf("%s/%s/models/%s:%s", info.ChannelBaseUrl, version, info.UpstreamModelName, action), nil
  112. }
  113. action := "generateContent"
  114. if info.IsStream {
  115. action = "streamGenerateContent?alt=sse"
  116. if info.RelayMode == constant.RelayModeGemini {
  117. info.DisablePing = true
  118. }
  119. }
  120. return fmt.Sprintf("%s/%s/models/%s:%s", info.ChannelBaseUrl, version, info.UpstreamModelName, action), nil
  121. }
  122. func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Header, info *relaycommon.RelayInfo) error {
  123. channel.SetupApiRequestHeader(info, c, req)
  124. req.Set("x-goog-api-key", info.ApiKey)
  125. return nil
  126. }
  127. func (a *Adaptor) ConvertOpenAIRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeneralOpenAIRequest) (any, error) {
  128. if request == nil {
  129. return nil, errors.New("request is nil")
  130. }
  131. geminiRequest, err := CovertGemini2OpenAI(c, *request, info)
  132. if err != nil {
  133. return nil, err
  134. }
  135. return geminiRequest, nil
  136. }
  137. func (a *Adaptor) ConvertRerankRequest(c *gin.Context, relayMode int, request dto.RerankRequest) (any, error) {
  138. return nil, nil
  139. }
  140. func (a *Adaptor) ConvertEmbeddingRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.EmbeddingRequest) (any, error) {
  141. if request.Input == nil {
  142. return nil, errors.New("input is required")
  143. }
  144. inputs := request.ParseInput()
  145. if len(inputs) == 0 {
  146. return nil, errors.New("input is empty")
  147. }
  148. // We always build a batch-style payload with `requests`, so ensure we call the
  149. // batch endpoint upstream to avoid payload/endpoint mismatches.
  150. info.IsGeminiBatchEmbedding = true
  151. // process all inputs
  152. geminiRequests := make([]map[string]interface{}, 0, len(inputs))
  153. for _, input := range inputs {
  154. geminiRequest := map[string]interface{}{
  155. "model": fmt.Sprintf("models/%s", info.UpstreamModelName),
  156. "content": dto.GeminiChatContent{
  157. Parts: []dto.GeminiPart{
  158. {
  159. Text: input,
  160. },
  161. },
  162. },
  163. }
  164. // set specific parameters for different models
  165. // https://ai.google.dev/api/embeddings?hl=zh-cn#method:-models.embedcontent
  166. switch info.UpstreamModelName {
  167. case "text-embedding-004", "gemini-embedding-exp-03-07", "gemini-embedding-001":
  168. // Only newer models introduced after 2024 support OutputDimensionality
  169. if request.Dimensions > 0 {
  170. geminiRequest["outputDimensionality"] = request.Dimensions
  171. }
  172. }
  173. geminiRequests = append(geminiRequests, geminiRequest)
  174. }
  175. return map[string]interface{}{
  176. "requests": geminiRequests,
  177. }, nil
  178. }
  179. func (a *Adaptor) ConvertOpenAIResponsesRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.OpenAIResponsesRequest) (any, error) {
  180. // TODO implement me
  181. return nil, errors.New("not implemented")
  182. }
  183. func (a *Adaptor) DoRequest(c *gin.Context, info *relaycommon.RelayInfo, requestBody io.Reader) (any, error) {
  184. return channel.DoApiRequest(a, c, info, requestBody)
  185. }
  186. func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, info *relaycommon.RelayInfo) (usage any, err *types.NewAPIError) {
  187. if info.RelayMode == constant.RelayModeGemini {
  188. if strings.HasSuffix(info.RequestURLPath, ":embedContent") ||
  189. strings.HasSuffix(info.RequestURLPath, ":batchEmbedContents") {
  190. return NativeGeminiEmbeddingHandler(c, resp, info)
  191. }
  192. if info.IsStream {
  193. return GeminiTextGenerationStreamHandler(c, info, resp)
  194. } else {
  195. return GeminiTextGenerationHandler(c, info, resp)
  196. }
  197. }
  198. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  199. return GeminiImageHandler(c, info, resp)
  200. }
  201. // check if the model is an embedding model
  202. if strings.HasPrefix(info.UpstreamModelName, "text-embedding") ||
  203. strings.HasPrefix(info.UpstreamModelName, "embedding") ||
  204. strings.HasPrefix(info.UpstreamModelName, "gemini-embedding") {
  205. return GeminiEmbeddingHandler(c, info, resp)
  206. }
  207. if info.IsStream {
  208. return GeminiChatStreamHandler(c, info, resp)
  209. } else {
  210. return GeminiChatHandler(c, info, resp)
  211. }
  212. }
  213. func (a *Adaptor) GetModelList() []string {
  214. return ModelList
  215. }
  216. func (a *Adaptor) GetChannelName() string {
  217. return ChannelName
  218. }