adaptor.go 7.4 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128129130131132133134135136137138139140141142143144145146147148149150151152153154155156157158159160161162163164165166167168169170171172173174175176177178179180181182183184185186187188189190191192193194195196197198199200201202203204205206207208209210211212213214215216217218219220221222223224225226227228229230231232233234235236237238239240241242243244245246
  1. package gemini
  2. import (
  3. "errors"
  4. "fmt"
  5. "io"
  6. "net/http"
  7. "one-api/dto"
  8. "one-api/relay/channel"
  9. "one-api/relay/channel/openai"
  10. relaycommon "one-api/relay/common"
  11. "one-api/relay/constant"
  12. "one-api/setting/model_setting"
  13. "one-api/types"
  14. "strings"
  15. "github.com/gin-gonic/gin"
  16. )
  17. type Adaptor struct {
  18. }
  19. func (a *Adaptor) ConvertGeminiRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeminiChatRequest) (any, error) {
  20. if len(request.Contents) > 0 {
  21. for i, content := range request.Contents {
  22. if i == 0 {
  23. if request.Contents[0].Role == "" {
  24. request.Contents[0].Role = "user"
  25. }
  26. }
  27. for _, part := range content.Parts {
  28. if part.FileData != nil {
  29. if part.FileData.MimeType == "" && strings.Contains(part.FileData.FileUri, "www.youtube.com") {
  30. part.FileData.MimeType = "video/webm"
  31. }
  32. }
  33. }
  34. }
  35. }
  36. return request, nil
  37. }
  38. func (a *Adaptor) ConvertClaudeRequest(c *gin.Context, info *relaycommon.RelayInfo, req *dto.ClaudeRequest) (any, error) {
  39. adaptor := openai.Adaptor{}
  40. oaiReq, err := adaptor.ConvertClaudeRequest(c, info, req)
  41. if err != nil {
  42. return nil, err
  43. }
  44. return a.ConvertOpenAIRequest(c, info, oaiReq.(*dto.GeneralOpenAIRequest))
  45. }
  46. func (a *Adaptor) ConvertAudioRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.AudioRequest) (io.Reader, error) {
  47. //TODO implement me
  48. return nil, errors.New("not implemented")
  49. }
  50. func (a *Adaptor) ConvertImageRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.ImageRequest) (any, error) {
  51. if !strings.HasPrefix(info.UpstreamModelName, "imagen") {
  52. return nil, errors.New("not supported model for image generation")
  53. }
  54. // convert size to aspect ratio
  55. aspectRatio := "1:1" // default aspect ratio
  56. switch request.Size {
  57. case "1024x1024":
  58. aspectRatio = "1:1"
  59. case "1024x1792":
  60. aspectRatio = "9:16"
  61. case "1792x1024":
  62. aspectRatio = "16:9"
  63. }
  64. // build gemini imagen request
  65. geminiRequest := dto.GeminiImageRequest{
  66. Instances: []dto.GeminiImageInstance{
  67. {
  68. Prompt: request.Prompt,
  69. },
  70. },
  71. Parameters: dto.GeminiImageParameters{
  72. SampleCount: request.N,
  73. AspectRatio: aspectRatio,
  74. PersonGeneration: "allow_adult", // default allow adult
  75. },
  76. }
  77. return geminiRequest, nil
  78. }
  79. func (a *Adaptor) Init(info *relaycommon.RelayInfo) {
  80. }
  81. func (a *Adaptor) GetRequestURL(info *relaycommon.RelayInfo) (string, error) {
  82. if model_setting.GetGeminiSettings().ThinkingAdapterEnabled {
  83. // 新增逻辑:处理 -thinking-<budget> 格式
  84. if strings.Contains(info.UpstreamModelName, "-thinking-") {
  85. parts := strings.Split(info.UpstreamModelName, "-thinking-")
  86. info.UpstreamModelName = parts[0]
  87. } else if strings.HasSuffix(info.UpstreamModelName, "-thinking") { // 旧的适配
  88. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-thinking")
  89. } else if strings.HasSuffix(info.UpstreamModelName, "-nothinking") {
  90. info.UpstreamModelName = strings.TrimSuffix(info.UpstreamModelName, "-nothinking")
  91. }
  92. }
  93. version := model_setting.GetGeminiVersionSetting(info.UpstreamModelName)
  94. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  95. return fmt.Sprintf("%s/%s/models/%s:predict", info.BaseUrl, version, info.UpstreamModelName), nil
  96. }
  97. if strings.HasPrefix(info.UpstreamModelName, "text-embedding") ||
  98. strings.HasPrefix(info.UpstreamModelName, "embedding") ||
  99. strings.HasPrefix(info.UpstreamModelName, "gemini-embedding") {
  100. return fmt.Sprintf("%s/%s/models/%s:batchEmbedContents", info.BaseUrl, version, info.UpstreamModelName), nil
  101. }
  102. action := "generateContent"
  103. if info.IsStream {
  104. action = "streamGenerateContent?alt=sse"
  105. }
  106. return fmt.Sprintf("%s/%s/models/%s:%s", info.BaseUrl, version, info.UpstreamModelName, action), nil
  107. }
  108. func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Header, info *relaycommon.RelayInfo) error {
  109. channel.SetupApiRequestHeader(info, c, req)
  110. req.Set("x-goog-api-key", info.ApiKey)
  111. return nil
  112. }
  113. func (a *Adaptor) ConvertOpenAIRequest(c *gin.Context, info *relaycommon.RelayInfo, request *dto.GeneralOpenAIRequest) (any, error) {
  114. if request == nil {
  115. return nil, errors.New("request is nil")
  116. }
  117. geminiRequest, err := CovertGemini2OpenAI(*request, info)
  118. if err != nil {
  119. return nil, err
  120. }
  121. return geminiRequest, nil
  122. }
  123. func (a *Adaptor) ConvertRerankRequest(c *gin.Context, relayMode int, request dto.RerankRequest) (any, error) {
  124. return nil, nil
  125. }
  126. func (a *Adaptor) ConvertEmbeddingRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.EmbeddingRequest) (any, error) {
  127. if request.Input == nil {
  128. return nil, errors.New("input is required")
  129. }
  130. inputs := request.ParseInput()
  131. if len(inputs) == 0 {
  132. return nil, errors.New("input is empty")
  133. }
  134. // process all inputs
  135. geminiRequests := make([]map[string]interface{}, 0, len(inputs))
  136. for _, input := range inputs {
  137. geminiRequest := map[string]interface{}{
  138. "model": fmt.Sprintf("models/%s", info.UpstreamModelName),
  139. "content": dto.GeminiChatContent{
  140. Parts: []dto.GeminiPart{
  141. {
  142. Text: input,
  143. },
  144. },
  145. },
  146. }
  147. // set specific parameters for different models
  148. // https://ai.google.dev/api/embeddings?hl=zh-cn#method:-models.embedcontent
  149. switch info.UpstreamModelName {
  150. case "text-embedding-004":
  151. // except embedding-001 supports setting `OutputDimensionality`
  152. if request.Dimensions > 0 {
  153. geminiRequest["outputDimensionality"] = request.Dimensions
  154. }
  155. }
  156. geminiRequests = append(geminiRequests, geminiRequest)
  157. }
  158. return map[string]interface{}{
  159. "requests": geminiRequests,
  160. }, nil
  161. }
  162. func (a *Adaptor) ConvertOpenAIResponsesRequest(c *gin.Context, info *relaycommon.RelayInfo, request dto.OpenAIResponsesRequest) (any, error) {
  163. // TODO implement me
  164. return nil, errors.New("not implemented")
  165. }
  166. func (a *Adaptor) DoRequest(c *gin.Context, info *relaycommon.RelayInfo, requestBody io.Reader) (any, error) {
  167. return channel.DoApiRequest(a, c, info, requestBody)
  168. }
  169. func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, info *relaycommon.RelayInfo) (usage any, err *types.NewAPIError) {
  170. if info.RelayMode == constant.RelayModeGemini {
  171. if info.IsStream {
  172. info.DisablePing = true
  173. return GeminiTextGenerationStreamHandler(c, info, resp)
  174. } else {
  175. return GeminiTextGenerationHandler(c, info, resp)
  176. }
  177. }
  178. if strings.HasPrefix(info.UpstreamModelName, "imagen") {
  179. return GeminiImageHandler(c, info, resp)
  180. }
  181. // check if the model is an embedding model
  182. if strings.HasPrefix(info.UpstreamModelName, "text-embedding") ||
  183. strings.HasPrefix(info.UpstreamModelName, "embedding") ||
  184. strings.HasPrefix(info.UpstreamModelName, "gemini-embedding") {
  185. return GeminiEmbeddingHandler(c, info, resp)
  186. }
  187. if info.IsStream {
  188. return GeminiChatStreamHandler(c, info, resp)
  189. } else {
  190. return GeminiChatHandler(c, info, resp)
  191. }
  192. //if usage.(*dto.Usage).CompletionTokenDetails.ReasoningTokens > 100 {
  193. // // 没有请求-thinking的情况下,产生思考token,则按照思考模型计费
  194. // if !strings.HasSuffix(info.OriginModelName, "-thinking") &&
  195. // !strings.HasSuffix(info.OriginModelName, "-nothinking") {
  196. // thinkingModelName := info.OriginModelName + "-thinking"
  197. // if operation_setting.SelfUseModeEnabled || helper.ContainPriceOrRatio(thinkingModelName) {
  198. // info.OriginModelName = thinkingModelName
  199. // }
  200. // }
  201. //}
  202. return nil, types.NewError(errors.New("not implemented"), types.ErrorCodeBadResponseBody)
  203. }
  204. func (a *Adaptor) GetModelList() []string {
  205. return ModelList
  206. }
  207. func (a *Adaptor) GetChannelName() string {
  208. return ChannelName
  209. }