Merge 009337ccf3 into 7ac553541b

feat: update openrouter models and price 20250213 (#2084 )
feat: update ali models and price 20250213 (#2086 )
2026-02-20 04:44:26 +08:00 · 2025-02-18 14:31:39 +08:00 · 2025-02-16 18:01:59 +08:00 · 2025-02-16 18:01:24 +08:00 · 2025-01-12 04:22:21 +00:00
32 changed files with 812 additions and 147 deletions
--- a/common/ctxkey/key.go
+++ b/common/ctxkey/key.go
@@ -21,4 +21,5 @@ const (
 	AvailableModels   = "available_models"
 	KeyRequestBody    = "key_request_body"
 	SystemPrompt      = "system_prompt"
+	Meta              = "meta"
 )
--- a/relay/adaptor/aiproxy/adaptor.go
+++ b/relay/adaptor/aiproxy/adaptor.go
@@ -38,7 +38,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return aiProxyLibraryRequest, nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/ali/adaptor.go
+++ b/relay/adaptor/ali/adaptor.go
@@ -67,7 +67,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	}
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/ali/constants.go
+++ b/relay/adaptor/ali/constants.go
@@ -14,10 +14,14 @@ var ModelList = []string{
 	"qwen2-72b-instruct", "qwen2-57b-a14b-instruct", "qwen2-7b-instruct", "qwen2-1.5b-instruct", "qwen2-0.5b-instruct",
 	"qwen1.5-110b-chat", "qwen1.5-72b-chat", "qwen1.5-32b-chat", "qwen1.5-14b-chat", "qwen1.5-7b-chat", "qwen1.5-1.8b-chat", "qwen1.5-0.5b-chat",
 	"qwen-72b-chat", "qwen-14b-chat", "qwen-7b-chat", "qwen-1.8b-chat", "qwen-1.8b-longcontext-chat",
+	"qvq-72b-preview",
+	"qwen2.5-vl-72b-instruct", "qwen2.5-vl-7b-instruct", "qwen2.5-vl-2b-instruct", "qwen2.5-vl-1b-instruct", "qwen2.5-vl-0.5b-instruct",
 	"qwen2-vl-7b-instruct", "qwen2-vl-2b-instruct", "qwen-vl-v1", "qwen-vl-chat-v1",
 	"qwen2-audio-instruct", "qwen-audio-chat",
 	"qwen2.5-math-72b-instruct", "qwen2.5-math-7b-instruct", "qwen2.5-math-1.5b-instruct", "qwen2-math-72b-instruct", "qwen2-math-7b-instruct", "qwen2-math-1.5b-instruct",
 	"qwen2.5-coder-32b-instruct", "qwen2.5-coder-14b-instruct", "qwen2.5-coder-7b-instruct", "qwen2.5-coder-3b-instruct", "qwen2.5-coder-1.5b-instruct", "qwen2.5-coder-0.5b-instruct",
 	"text-embedding-v1", "text-embedding-v3", "text-embedding-v2", "text-embedding-async-v2", "text-embedding-async-v1",
 	"ali-stable-diffusion-xl", "ali-stable-diffusion-v1.5", "wanx-v1",
+	"qwen-mt-plus", "qwen-mt-turbo",
+	"deepseek-r1", "deepseek-v3", "deepseek-r1-distill-qwen-1.5b", "deepseek-r1-distill-qwen-7b", "deepseek-r1-distill-qwen-14b", "deepseek-r1-distill-qwen-32b", "deepseek-r1-distill-llama-8b", "deepseek-r1-distill-llama-70b",
 }
--- a/relay/adaptor/anthropic/adaptor.go
+++ b/relay/adaptor/anthropic/adaptor.go
@@ -50,7 +50,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return ConvertRequest(*request), nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/aws/adaptor.go
+++ b/relay/adaptor/aws/adaptor.go
@@ -72,7 +72,7 @@ func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Request, meta *me
 	return nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/aws/utils/adaptor.go
+++ b/relay/adaptor/aws/utils/adaptor.go
@@ -39,7 +39,7 @@ func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Request, meta *me
 	return nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/baidu/adaptor.go
+++ b/relay/adaptor/baidu/adaptor.go
@@ -109,7 +109,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	}
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/cloudflare/adaptor.go
+++ b/relay/adaptor/cloudflare/adaptor.go
@@ -19,7 +19,7 @@ type Adaptor struct {
 }

 // ConvertImageRequest implements adaptor.Adaptor.
-func (*Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (*Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	return nil, errors.New("not implemented")
 }

--- a/relay/adaptor/cohere/adaptor.go
+++ b/relay/adaptor/cohere/adaptor.go
@@ -15,7 +15,7 @@ import (
 type Adaptor struct{}

 // ConvertImageRequest implements adaptor.Adaptor.
-func (*Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (*Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	return nil, errors.New("not implemented")
 }

--- a/relay/adaptor/coze/adaptor.go
+++ b/relay/adaptor/coze/adaptor.go
@@ -38,7 +38,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return ConvertRequest(*request), nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/deepl/adaptor.go
+++ b/relay/adaptor/deepl/adaptor.go
@@ -39,7 +39,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return convertedRequest, nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/gemini/adaptor.go
+++ b/relay/adaptor/gemini/adaptor.go
@@ -66,7 +66,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	}
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/interface.go
+++ b/relay/adaptor/interface.go
@@ -13,7 +13,7 @@ type Adaptor interface {
 	GetRequestURL(meta *meta.Meta) (string, error)
 	SetupRequestHeader(c *gin.Context, req *http.Request, meta *meta.Meta) error
 	ConvertRequest(c *gin.Context, relayMode int, request *model.GeneralOpenAIRequest) (any, error)
-	ConvertImageRequest(request *model.ImageRequest) (any, error)
+	ConvertImageRequest(c *gin.Context, request *model.ImageRequest) (any, error)
 	DoRequest(c *gin.Context, meta *meta.Meta, requestBody io.Reader) (*http.Response, error)
 	DoResponse(c *gin.Context, resp *http.Response, meta *meta.Meta) (usage *model.Usage, err *model.ErrorWithStatusCode)
 	GetModelList() []string
--- a/relay/adaptor/ollama/adaptor.go
+++ b/relay/adaptor/ollama/adaptor.go
@@ -48,7 +48,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	}
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/openai/adaptor.go
+++ b/relay/adaptor/openai/adaptor.go
@@ -95,7 +95,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return request, nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/openrouter/constants.go
+++ b/relay/adaptor/openrouter/constants.go
@@ -1,20 +1,235 @@
 package openrouter

 var ModelList = []string{
-	"openai/gpt-3.5-turbo",
-	"openai/chatgpt-4o-latest",
-	"openai/o1",
-	"openai/o1-preview",
-	"openai/o1-mini",
-	"openai/o3-mini",
-	"google/gemini-2.0-flash-001",
-	"google/gemini-2.0-flash-thinking-exp:free",
-	"google/gemini-2.0-flash-lite-preview-02-05:free",
-	"google/gemini-2.0-pro-exp-02-05:free",
-	"google/gemini-flash-1.5-8b",
-	"anthropic/claude-3.5-sonnet",
+	"01-ai/yi-large",
+	"aetherwiing/mn-starcannon-12b",
+	"ai21/jamba-1-5-large",
+	"ai21/jamba-1-5-mini",
+	"ai21/jamba-instruct",
+	"aion-labs/aion-1.0",
+	"aion-labs/aion-1.0-mini",
+	"aion-labs/aion-rp-llama-3.1-8b",
+	"allenai/llama-3.1-tulu-3-405b",
+	"alpindale/goliath-120b",
+	"alpindale/magnum-72b",
+	"amazon/nova-lite-v1",
+	"amazon/nova-micro-v1",
+	"amazon/nova-pro-v1",
+	"anthracite-org/magnum-v2-72b",
+	"anthracite-org/magnum-v4-72b",
+	"anthropic/claude-2",
+	"anthropic/claude-2.0",
+	"anthropic/claude-2.0:beta",
+	"anthropic/claude-2.1",
+	"anthropic/claude-2.1:beta",
+	"anthropic/claude-2:beta",
+	"anthropic/claude-3-haiku",
+	"anthropic/claude-3-haiku:beta",
+	"anthropic/claude-3-opus",
+	"anthropic/claude-3-opus:beta",
+	"anthropic/claude-3-sonnet",
+	"anthropic/claude-3-sonnet:beta",
 	"anthropic/claude-3.5-haiku",
-	"deepseek/deepseek-r1:free",
+	"anthropic/claude-3.5-haiku-20241022",
+	"anthropic/claude-3.5-haiku-20241022:beta",
+	"anthropic/claude-3.5-haiku:beta",
+	"anthropic/claude-3.5-sonnet",
+	"anthropic/claude-3.5-sonnet-20240620",
+	"anthropic/claude-3.5-sonnet-20240620:beta",
+	"anthropic/claude-3.5-sonnet:beta",
+	"cognitivecomputations/dolphin-mixtral-8x22b",
+	"cognitivecomputations/dolphin-mixtral-8x7b",
+	"cohere/command",
+	"cohere/command-r",
+	"cohere/command-r-03-2024",
+	"cohere/command-r-08-2024",
+	"cohere/command-r-plus",
+	"cohere/command-r-plus-04-2024",
+	"cohere/command-r-plus-08-2024",
+	"cohere/command-r7b-12-2024",
+	"databricks/dbrx-instruct",
+	"deepseek/deepseek-chat",
+	"deepseek/deepseek-chat-v2.5",
+	"deepseek/deepseek-chat:free",
 	"deepseek/deepseek-r1",
+	"deepseek/deepseek-r1-distill-llama-70b",
+	"deepseek/deepseek-r1-distill-llama-70b:free",
+	"deepseek/deepseek-r1-distill-llama-8b",
+	"deepseek/deepseek-r1-distill-qwen-1.5b",
+	"deepseek/deepseek-r1-distill-qwen-14b",
+	"deepseek/deepseek-r1-distill-qwen-32b",
+	"deepseek/deepseek-r1:free",
+	"eva-unit-01/eva-llama-3.33-70b",
+	"eva-unit-01/eva-qwen-2.5-32b",
+	"eva-unit-01/eva-qwen-2.5-72b",
+	"google/gemini-2.0-flash-001",
+	"google/gemini-2.0-flash-exp:free",
+	"google/gemini-2.0-flash-lite-preview-02-05:free",
+	"google/gemini-2.0-flash-thinking-exp-1219:free",
+	"google/gemini-2.0-flash-thinking-exp:free",
+	"google/gemini-2.0-pro-exp-02-05:free",
+	"google/gemini-exp-1206:free",
+	"google/gemini-flash-1.5",
+	"google/gemini-flash-1.5-8b",
+	"google/gemini-flash-1.5-8b-exp",
+	"google/gemini-pro",
+	"google/gemini-pro-1.5",
+	"google/gemini-pro-vision",
+	"google/gemma-2-27b-it",
+	"google/gemma-2-9b-it",
+	"google/gemma-2-9b-it:free",
+	"google/gemma-7b-it",
+	"google/learnlm-1.5-pro-experimental:free",
+	"google/palm-2-chat-bison",
+	"google/palm-2-chat-bison-32k",
+	"google/palm-2-codechat-bison",
+	"google/palm-2-codechat-bison-32k",
+	"gryphe/mythomax-l2-13b",
+	"gryphe/mythomax-l2-13b:free",
+	"huggingfaceh4/zephyr-7b-beta:free",
+	"infermatic/mn-inferor-12b",
+	"inflection/inflection-3-pi",
+	"inflection/inflection-3-productivity",
+	"jondurbin/airoboros-l2-70b",
+	"liquid/lfm-3b",
+	"liquid/lfm-40b",
+	"liquid/lfm-7b",
+	"mancer/weaver",
+	"meta-llama/llama-2-13b-chat",
+	"meta-llama/llama-2-70b-chat",
+	"meta-llama/llama-3-70b-instruct",
+	"meta-llama/llama-3-8b-instruct",
+	"meta-llama/llama-3-8b-instruct:free",
+	"meta-llama/llama-3.1-405b",
+	"meta-llama/llama-3.1-405b-instruct",
+	"meta-llama/llama-3.1-70b-instruct",
+	"meta-llama/llama-3.1-8b-instruct",
+	"meta-llama/llama-3.2-11b-vision-instruct",
+	"meta-llama/llama-3.2-11b-vision-instruct:free",
+	"meta-llama/llama-3.2-1b-instruct",
+	"meta-llama/llama-3.2-3b-instruct",
+	"meta-llama/llama-3.2-90b-vision-instruct",
+	"meta-llama/llama-3.3-70b-instruct",
+	"meta-llama/llama-3.3-70b-instruct:free",
+	"meta-llama/llama-guard-2-8b",
+	"microsoft/phi-3-medium-128k-instruct",
+	"microsoft/phi-3-medium-128k-instruct:free",
+	"microsoft/phi-3-mini-128k-instruct",
+	"microsoft/phi-3-mini-128k-instruct:free",
+	"microsoft/phi-3.5-mini-128k-instruct",
+	"microsoft/phi-4",
+	"microsoft/wizardlm-2-7b",
+	"microsoft/wizardlm-2-8x22b",
+	"minimax/minimax-01",
+	"mistralai/codestral-2501",
+	"mistralai/codestral-mamba",
+	"mistralai/ministral-3b",
+	"mistralai/ministral-8b",
+	"mistralai/mistral-7b-instruct",
+	"mistralai/mistral-7b-instruct-v0.1",
+	"mistralai/mistral-7b-instruct-v0.3",
+	"mistralai/mistral-7b-instruct:free",
+	"mistralai/mistral-large",
+	"mistralai/mistral-large-2407",
+	"mistralai/mistral-large-2411",
+	"mistralai/mistral-medium",
+	"mistralai/mistral-nemo",
+	"mistralai/mistral-nemo:free",
+	"mistralai/mistral-small",
+	"mistralai/mistral-small-24b-instruct-2501",
+	"mistralai/mistral-small-24b-instruct-2501:free",
+	"mistralai/mistral-tiny",
+	"mistralai/mixtral-8x22b-instruct",
+	"mistralai/mixtral-8x7b",
+	"mistralai/mixtral-8x7b-instruct",
+	"mistralai/pixtral-12b",
+	"mistralai/pixtral-large-2411",
+	"neversleep/llama-3-lumimaid-70b",
+	"neversleep/llama-3-lumimaid-8b",
+	"neversleep/llama-3-lumimaid-8b:extended",
+	"neversleep/llama-3.1-lumimaid-70b",
+	"neversleep/llama-3.1-lumimaid-8b",
+	"neversleep/noromaid-20b",
+	"nothingiisreal/mn-celeste-12b",
+	"nousresearch/hermes-2-pro-llama-3-8b",
+	"nousresearch/hermes-3-llama-3.1-405b",
+	"nousresearch/hermes-3-llama-3.1-70b",
+	"nousresearch/nous-hermes-2-mixtral-8x7b-dpo",
+	"nousresearch/nous-hermes-llama2-13b",
+	"nvidia/llama-3.1-nemotron-70b-instruct",
+	"nvidia/llama-3.1-nemotron-70b-instruct:free",
+	"openai/chatgpt-4o-latest",
+	"openai/gpt-3.5-turbo",
+	"openai/gpt-3.5-turbo-0125",
+	"openai/gpt-3.5-turbo-0613",
+	"openai/gpt-3.5-turbo-1106",
+	"openai/gpt-3.5-turbo-16k",
+	"openai/gpt-3.5-turbo-instruct",
+	"openai/gpt-4",
+	"openai/gpt-4-0314",
+	"openai/gpt-4-1106-preview",
+	"openai/gpt-4-32k",
+	"openai/gpt-4-32k-0314",
+	"openai/gpt-4-turbo",
+	"openai/gpt-4-turbo-preview",
+	"openai/gpt-4o",
+	"openai/gpt-4o-2024-05-13",
+	"openai/gpt-4o-2024-08-06",
+	"openai/gpt-4o-2024-11-20",
+	"openai/gpt-4o-mini",
+	"openai/gpt-4o-mini-2024-07-18",
+	"openai/gpt-4o:extended",
+	"openai/o1",
+	"openai/o1-mini",
+	"openai/o1-mini-2024-09-12",
+	"openai/o1-preview",
+	"openai/o1-preview-2024-09-12",
+	"openai/o3-mini",
+	"openai/o3-mini-high",
+	"openchat/openchat-7b",
+	"openchat/openchat-7b:free",
+	"openrouter/auto",
+	"perplexity/llama-3.1-sonar-huge-128k-online",
+	"perplexity/llama-3.1-sonar-large-128k-chat",
+	"perplexity/llama-3.1-sonar-large-128k-online",
+	"perplexity/llama-3.1-sonar-small-128k-chat",
+	"perplexity/llama-3.1-sonar-small-128k-online",
+	"perplexity/sonar",
+	"perplexity/sonar-reasoning",
+	"pygmalionai/mythalion-13b",
+	"qwen/qvq-72b-preview",
+	"qwen/qwen-2-72b-instruct",
+	"qwen/qwen-2-7b-instruct",
+	"qwen/qwen-2-7b-instruct:free",
+	"qwen/qwen-2-vl-72b-instruct",
+	"qwen/qwen-2-vl-7b-instruct",
+	"qwen/qwen-2.5-72b-instruct",
+	"qwen/qwen-2.5-7b-instruct",
+	"qwen/qwen-2.5-coder-32b-instruct",
+	"qwen/qwen-max",
+	"qwen/qwen-plus",
+	"qwen/qwen-turbo",
 	"qwen/qwen-vl-plus:free",
+	"qwen/qwen2.5-vl-72b-instruct:free",
+	"qwen/qwq-32b-preview",
+	"raifle/sorcererlm-8x22b",
+	"sao10k/fimbulvetr-11b-v2",
+	"sao10k/l3-euryale-70b",
+	"sao10k/l3-lunaris-8b",
+	"sao10k/l3.1-70b-hanami-x1",
+	"sao10k/l3.1-euryale-70b",
+	"sao10k/l3.3-euryale-70b",
+	"sophosympatheia/midnight-rose-70b",
+	"sophosympatheia/rogue-rose-103b-v0.2:free",
+	"teknium/openhermes-2.5-mistral-7b",
+	"thedrummer/rocinante-12b",
+	"thedrummer/unslopnemo-12b",
+	"undi95/remm-slerp-l2-13b",
+	"undi95/toppy-m-7b",
+	"undi95/toppy-m-7b:free",
+	"x-ai/grok-2-1212",
+	"x-ai/grok-2-vision-1212",
+	"x-ai/grok-beta",
+	"x-ai/grok-vision-beta",
+	"xwin-lm/xwin-lm-70b",
 }
--- a/relay/adaptor/palm/adaptor.go
+++ b/relay/adaptor/palm/adaptor.go
@@ -36,7 +36,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return ConvertRequest(*request), nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/proxy/adaptor.go
+++ b/relay/adaptor/proxy/adaptor.go
@@ -80,7 +80,7 @@ func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Request, meta *me
 	return nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	return nil, errors.Errorf("not implement")
 }

--- a/relay/adaptor/replicate/adaptor.go
+++ b/relay/adaptor/replicate/adaptor.go
@@ -23,7 +23,29 @@ type Adaptor struct {
 }

 // ConvertImageRequest implements adaptor.Adaptor.
-func (*Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
+	return nil, errors.New("should call replicate.ConvertImageRequest instead")
+}
+
+func ConvertImageRequest(c *gin.Context, request *model.ImageRequest) (any, error) {
+	meta := meta.GetByContext(c)
+
+	if request.ResponseFormat != "b64_json" {
+		return nil, errors.New("only support b64_json response format")
+	}
+	if request.N != 1 && request.N != 0 {
+		return nil, errors.New("only support N=1")
+	}
+
+	switch meta.Mode {
+	case relaymode.ImagesGenerations:
+		return convertImageCreateRequest(request)
+	default:
+		return nil, errors.New("not implemented")
+	}
+}
+
+func convertImageCreateRequest(request *model.ImageRequest) (any, error) {
 	return DrawImageRequest{
 		Input: ImageInput{
 			Steps:           25,
--- a/relay/adaptor/tencent/adaptor.go
+++ b/relay/adaptor/tencent/adaptor.go
@@ -69,7 +69,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return convertedRequest, nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/vertexai/adaptor.go
+++ b/relay/adaptor/vertexai/adaptor.go
@@ -1,18 +1,20 @@
 package vertexai

 import (
-	"errors"
 	"fmt"
 	"io"
 	"net/http"
+	"slices"
 	"strings"

 	"github.com/gin-gonic/gin"
+	"github.com/pkg/errors"
 	"github.com/songquanpeng/one-api/relay/adaptor"
 	channelhelper "github.com/songquanpeng/one-api/relay/adaptor"
+	"github.com/songquanpeng/one-api/relay/adaptor/vertexai/imagen"
 	"github.com/songquanpeng/one-api/relay/meta"
 	"github.com/songquanpeng/one-api/relay/model"
-	relaymodel "github.com/songquanpeng/one-api/relay/model"
+	relayModel "github.com/songquanpeng/one-api/relay/model"
 )

 var _ adaptor.Adaptor = new(Adaptor)
@@ -24,14 +26,27 @@ type Adaptor struct{}
 func (a *Adaptor) Init(meta *meta.Meta) {
 }

-func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.GeneralOpenAIRequest) (any, error) {
-	if request == nil {
-		return nil, errors.New("request is nil")
+func (a *Adaptor) ConvertImageRequest(c *gin.Context, request *model.ImageRequest) (any, error) {
+	meta := meta.GetByContext(c)
+
+	if request.ResponseFormat != "b64_json" {
+		return nil, errors.New("only support b64_json response format")
 	}

-	adaptor := GetAdaptor(request.Model)
+	adaptor := GetAdaptor(meta.ActualModelName)
 	if adaptor == nil {
-		return nil, errors.New("adaptor not found")
+		return nil, errors.Errorf("cannot found vertex image adaptor for model %s", meta.ActualModelName)
+	}
+
+	return adaptor.ConvertImageRequest(c, request)
+}
+
+func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.GeneralOpenAIRequest) (any, error) {
+	meta := meta.GetByContext(c)
+
+	adaptor := GetAdaptor(meta.ActualModelName)
+	if adaptor == nil {
+		return nil, errors.Errorf("cannot found vertex chat adaptor for model %s", meta.ActualModelName)
 	}

 	return adaptor.ConvertRequest(c, relayMode, request)
@@ -40,9 +55,9 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, meta *meta.Meta) (usage *model.Usage, err *model.ErrorWithStatusCode) {
 	adaptor := GetAdaptor(meta.ActualModelName)
 	if adaptor == nil {
-		return nil, &relaymodel.ErrorWithStatusCode{
+		return nil, &relayModel.ErrorWithStatusCode{
 			StatusCode: http.StatusInternalServerError,
-			Error: relaymodel.Error{
+			Error: relayModel.Error{
 				Message: "adaptor not found",
 			},
 		}
@@ -60,14 +75,19 @@ func (a *Adaptor) GetChannelName() string {
 }

 func (a *Adaptor) GetRequestURL(meta *meta.Meta) (string, error) {
-	suffix := ""
-	if strings.HasPrefix(meta.ActualModelName, "gemini") {
+	var suffix string
+	switch {
+	case strings.HasPrefix(meta.ActualModelName, "gemini"):
 		if meta.IsStream {
 			suffix = "streamGenerateContent?alt=sse"
 		} else {
 			suffix = "generateContent"
 		}
-	} else {
+	case slices.Contains(imagen.ModelList, meta.ActualModelName):
+		return fmt.Sprintf("https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/google/models/imagen-3.0-generate-001:predict",
+			meta.Config.Region, meta.Config.VertexAIProjectID, meta.Config.Region,
+		), nil
+	default:
 		if meta.IsStream {
 			suffix = "streamRawPredict?alt=sse"
 		} else {
@@ -85,6 +105,7 @@ func (a *Adaptor) GetRequestURL(meta *meta.Meta) (string, error) {
 			suffix,
 		), nil
 	}
+
 	return fmt.Sprintf(
 		"https://%s-aiplatform.googleapis.com/v1/projects/%s/locations/%s/publishers/google/models/%s:%s",
 		meta.Config.Region,
@@ -105,13 +126,6 @@ func (a *Adaptor) SetupRequestHeader(c *gin.Context, req *http.Request, meta *me
 	return nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
-	if request == nil {
-		return nil, errors.New("request is nil")
-	}
-	return request, nil
-}
-
 func (a *Adaptor) DoRequest(c *gin.Context, meta *meta.Meta, requestBody io.Reader) (*http.Response, error) {
 	return channelhelper.DoRequestHelper(a, c, meta, requestBody)
 }
--- a/relay/adaptor/vertexai/claude/adapter.go
+++ b/relay/adaptor/vertexai/claude/adapter.go
@@ -50,6 +50,10 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return req, nil
 }

+func (a *Adaptor) ConvertImageRequest(c *gin.Context, request *model.ImageRequest) (any, error) {
+	return nil, errors.New("not support image request")
+}
+
 func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, meta *meta.Meta) (usage *model.Usage, err *model.ErrorWithStatusCode) {
 	if meta.IsStream {
 		err, usage = anthropic.StreamHandler(c, resp)
--- a/relay/adaptor/vertexai/gemini/adapter.go
+++ b/relay/adaptor/vertexai/gemini/adapter.go
@@ -38,6 +38,10 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return geminiRequest, nil
 }

+func (a *Adaptor) ConvertImageRequest(c *gin.Context, request *model.ImageRequest) (any, error) {
+	return nil, errors.New("not support image request")
+}
+
 func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, meta *meta.Meta) (usage *model.Usage, err *model.ErrorWithStatusCode) {
 	if meta.IsStream {
 		var responseText string
--- a/relay/adaptor/vertexai/imagen/adaptor.go
+++ b/relay/adaptor/vertexai/imagen/adaptor.go
@@ -0,0 +1,105 @@
+package imagen
+
+import (
+	"encoding/json"
+	"io"
+	"net/http"
+	"time"
+
+	"github.com/gin-gonic/gin"
+	"github.com/pkg/errors"
+	"github.com/songquanpeng/one-api/relay/adaptor/openai"
+	"github.com/songquanpeng/one-api/relay/meta"
+	"github.com/songquanpeng/one-api/relay/model"
+	"github.com/songquanpeng/one-api/relay/relaymode"
+)
+
+var ModelList = []string{
+	"imagen-3.0-generate-001",
+}
+
+type Adaptor struct {
+}
+
+func (a *Adaptor) ConvertImageRequest(c *gin.Context, request *model.ImageRequest) (any, error) {
+	meta := meta.GetByContext(c)
+
+	if request.ResponseFormat != "b64_json" {
+		return nil, errors.New("only support b64_json response format")
+	}
+	if request.N <= 0 {
+		return nil, errors.New("n must be greater than 0")
+	}
+
+	switch meta.Mode {
+	case relaymode.ImagesGenerations:
+		return convertImageCreateRequest(request)
+	default:
+		return nil, errors.New("not implemented")
+	}
+}
+
+func convertImageCreateRequest(request *model.ImageRequest) (any, error) {
+	return CreateImageRequest{
+		Instances: []createImageInstance{
+			{
+				Prompt: request.Prompt,
+			},
+		},
+		Parameters: createImageParameters{
+			SampleCount: request.N,
+		},
+	}, nil
+}
+
+func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.GeneralOpenAIRequest) (any, error) {
+	return nil, errors.New("not implemented")
+}
+
+func (a *Adaptor) DoResponse(c *gin.Context, resp *http.Response, meta *meta.Meta) (usage *model.Usage, wrapErr *model.ErrorWithStatusCode) {
+	respBody, err := io.ReadAll(resp.Body)
+	if err != nil {
+		return nil, openai.ErrorWrapper(
+			errors.Wrap(err, "failed to read response body"),
+			"read_response_body",
+			http.StatusInternalServerError,
+		)
+	}
+
+	if resp.StatusCode != http.StatusOK {
+		return nil, openai.ErrorWrapper(
+			errors.Errorf("upstream response status code: %d, body: %s", resp.StatusCode, string(respBody)),
+			"upstream_response",
+			http.StatusInternalServerError,
+		)
+	}
+
+	imagenResp := new(CreateImageResponse)
+	if err := json.Unmarshal(respBody, imagenResp); err != nil {
+		return nil, openai.ErrorWrapper(
+			errors.Wrap(err, "failed to decode response body"),
+			"unmarshal_upstream_response",
+			http.StatusInternalServerError,
+		)
+	}
+
+	if len(imagenResp.Predictions) == 0 {
+		return nil, openai.ErrorWrapper(
+			errors.New("empty predictions"),
+			"empty_predictions",
+			http.StatusInternalServerError,
+		)
+	}
+
+	oaiResp := openai.ImageResponse{
+		Created: time.Now().Unix(),
+	}
+	for _, prediction := range imagenResp.Predictions {
+		oaiResp.Data = append(oaiResp.Data, openai.ImageData{
+			B64Json: prediction.BytesBase64Encoded,
+		})
+	}
+
+	c.JSON(http.StatusOK, oaiResp)
+	return nil, nil
+}
--- a/relay/adaptor/vertexai/imagen/model.go
+++ b/relay/adaptor/vertexai/imagen/model.go
@@ -0,0 +1,23 @@
+package imagen
+
+type CreateImageRequest struct {
+	Instances  []createImageInstance `json:"instances" binding:"required,min=1"`
+	Parameters createImageParameters `json:"parameters" binding:"required"`
+}
+
+type createImageInstance struct {
+	Prompt string `json:"prompt"`
+}
+
+type createImageParameters struct {
+	SampleCount int `json:"sample_count" binding:"required,min=1"`
+}
+
+type CreateImageResponse struct {
+	Predictions []createImageResponsePrediction `json:"predictions"`
+}
+
+type createImageResponsePrediction struct {
+	MimeType           string `json:"mimeType"`
+	BytesBase64Encoded string `json:"bytesBase64Encoded"`
+}
--- a/relay/adaptor/vertexai/registry.go
+++ b/relay/adaptor/vertexai/registry.go
@@ -6,6 +6,7 @@ import (
 	"github.com/gin-gonic/gin"
 	claude "github.com/songquanpeng/one-api/relay/adaptor/vertexai/claude"
 	gemini "github.com/songquanpeng/one-api/relay/adaptor/vertexai/gemini"
+	"github.com/songquanpeng/one-api/relay/adaptor/vertexai/imagen"
 	"github.com/songquanpeng/one-api/relay/meta"
 	"github.com/songquanpeng/one-api/relay/model"
 )
@@ -13,8 +14,9 @@ import (
 type VertexAIModelType int

 const (
-	VerterAIClaude VertexAIModelType = iota + 1
-	VerterAIGemini
+	VertexAIClaude VertexAIModelType = iota + 1
+	VertexAIGemini
+	VertexAIImagen
 )

 var modelMapping = map[string]VertexAIModelType{}
@@ -23,27 +25,35 @@ var modelList = []string{}
 func init() {
 	modelList = append(modelList, claude.ModelList...)
 	for _, model := range claude.ModelList {
-		modelMapping[model] = VerterAIClaude
+		modelMapping[model] = VertexAIClaude
 	}

 	modelList = append(modelList, gemini.ModelList...)
 	for _, model := range gemini.ModelList {
-		modelMapping[model] = VerterAIGemini
+		modelMapping[model] = VertexAIGemini
+	}
+
+	modelList = append(modelList, imagen.ModelList...)
+	for _, model := range imagen.ModelList {
+		modelMapping[model] = VertexAIImagen
 	}
 }

 type innerAIAdapter interface {
 	ConvertRequest(c *gin.Context, relayMode int, request *model.GeneralOpenAIRequest) (any, error)
+	ConvertImageRequest(c *gin.Context, request *model.ImageRequest) (any, error)
 	DoResponse(c *gin.Context, resp *http.Response, meta *meta.Meta) (usage *model.Usage, err *model.ErrorWithStatusCode)
 }

 func GetAdaptor(model string) innerAIAdapter {
 	adaptorType := modelMapping[model]
 	switch adaptorType {
-	case VerterAIClaude:
+	case VertexAIClaude:
 		return &claude.Adaptor{}
-	case VerterAIGemini:
+	case VertexAIGemini:
 		return &gemini.Adaptor{}
+	case VertexAIImagen:
+		return &imagen.Adaptor{}
 	default:
 		return nil
 	}
--- a/relay/adaptor/xunfei/adaptor.go
+++ b/relay/adaptor/xunfei/adaptor.go
@@ -39,7 +39,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	return nil, nil
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/adaptor/zhipu/adaptor.go
+++ b/relay/adaptor/zhipu/adaptor.go
@@ -80,7 +80,7 @@ func (a *Adaptor) ConvertRequest(c *gin.Context, relayMode int, request *model.G
 	}
 }

-func (a *Adaptor) ConvertImageRequest(request *model.ImageRequest) (any, error) {
+func (a *Adaptor) ConvertImageRequest(_ *gin.Context, request *model.ImageRequest) (any, error) {
 	if request == nil {
 		return nil, errors.New("request is nil")
 	}
--- a/relay/billing/ratio/model.go
+++ b/relay/billing/ratio/model.go
@@ -59,6 +59,8 @@ var ModelRatio = map[string]float64{
 	"o1-preview-2024-09-12":   7.5,
 	"o1-mini":                 1.5, // $3.00 / 1M input tokens
 	"o1-mini-2024-09-12":      1.5,
+	"o3-mini":                 1.5, // $3.00 / 1M input tokens
+	"o3-mini-2025-01-31":      1.5,
 	"davinci-002":             1,   // $0.002 / 1K tokens
 	"babbage-002":             0.2, // $0.0004 / 1K tokens
 	"text-ada-001":            0.2,
@@ -159,91 +161,105 @@ var ModelRatio = map[string]float64{
 	"embedding-2":      0.0005 * RMB,
 	"embedding-3":      0.0005 * RMB,
 	// https://help.aliyun.com/zh/dashscope/developer-reference/tongyi-thousand-questions-metering-and-billing
-	"qwen-turbo":                  1.4286, // ￥0.02 / 1k tokens
-	"qwen-turbo-latest":           1.4286,
-	"qwen-plus":                   1.4286,
-	"qwen-plus-latest":            1.4286,
-	"qwen-max":                    1.4286,
-	"qwen-max-latest":             1.4286,
-	"qwen-max-longcontext":        1.4286,
-	"qwen-vl-max":                 1.4286,
-	"qwen-vl-max-latest":          1.4286,
-	"qwen-vl-plus":                1.4286,
-	"qwen-vl-plus-latest":         1.4286,
-	"qwen-vl-ocr":                 1.4286,
-	"qwen-vl-ocr-latest":          1.4286,
-	"qwen-audio-turbo":            1.4286,
-	"qwen-math-plus":              1.4286,
-	"qwen-math-plus-latest":       1.4286,
-	"qwen-math-turbo":             1.4286,
-	"qwen-math-turbo-latest":      1.4286,
-	"qwen-coder-plus":             1.4286,
-	"qwen-coder-plus-latest":      1.4286,
-	"qwen-coder-turbo":            1.4286,
-	"qwen-coder-turbo-latest":     1.4286,
-	"qwq-32b-preview":             1.4286,
-	"qwen2.5-72b-instruct":        1.4286,
-	"qwen2.5-32b-instruct":        1.4286,
-	"qwen2.5-14b-instruct":        1.4286,
-	"qwen2.5-7b-instruct":         1.4286,
-	"qwen2.5-3b-instruct":         1.4286,
-	"qwen2.5-1.5b-instruct":       1.4286,
-	"qwen2.5-0.5b-instruct":       1.4286,
-	"qwen2-72b-instruct":          1.4286,
-	"qwen2-57b-a14b-instruct":     1.4286,
-	"qwen2-7b-instruct":           1.4286,
-	"qwen2-1.5b-instruct":         1.4286,
-	"qwen2-0.5b-instruct":         1.4286,
-	"qwen1.5-110b-chat":           1.4286,
-	"qwen1.5-72b-chat":            1.4286,
-	"qwen1.5-32b-chat":            1.4286,
-	"qwen1.5-14b-chat":            1.4286,
-	"qwen1.5-7b-chat":             1.4286,
-	"qwen1.5-1.8b-chat":           1.4286,
-	"qwen1.5-0.5b-chat":           1.4286,
-	"qwen-72b-chat":               1.4286,
-	"qwen-14b-chat":               1.4286,
-	"qwen-7b-chat":                1.4286,
-	"qwen-1.8b-chat":              1.4286,
-	"qwen-1.8b-longcontext-chat":  1.4286,
-	"qwen2-vl-7b-instruct":        1.4286,
-	"qwen2-vl-2b-instruct":        1.4286,
-	"qwen-vl-v1":                  1.4286,
-	"qwen-vl-chat-v1":             1.4286,
-	"qwen2-audio-instruct":        1.4286,
-	"qwen-audio-chat":             1.4286,
-	"qwen2.5-math-72b-instruct":   1.4286,
-	"qwen2.5-math-7b-instruct":    1.4286,
-	"qwen2.5-math-1.5b-instruct":  1.4286,
-	"qwen2-math-72b-instruct":     1.4286,
-	"qwen2-math-7b-instruct":      1.4286,
-	"qwen2-math-1.5b-instruct":    1.4286,
-	"qwen2.5-coder-32b-instruct":  1.4286,
-	"qwen2.5-coder-14b-instruct":  1.4286,
-	"qwen2.5-coder-7b-instruct":   1.4286,
-	"qwen2.5-coder-3b-instruct":   1.4286,
-	"qwen2.5-coder-1.5b-instruct": 1.4286,
-	"qwen2.5-coder-0.5b-instruct": 1.4286,
-	"text-embedding-v1":           0.05, // ￥0.0007 / 1k tokens
-	"text-embedding-v3":           0.05,
-	"text-embedding-v2":           0.05,
-	"text-embedding-async-v2":     0.05,
-	"text-embedding-async-v1":     0.05,
-	"ali-stable-diffusion-xl":     8.00,
-	"ali-stable-diffusion-v1.5":   8.00,
-	"wanx-v1":                     8.00,
-	"SparkDesk":                   1.2858, // ￥0.018 / 1k tokens
-	"SparkDesk-v1.1":              1.2858, // ￥0.018 / 1k tokens
-	"SparkDesk-v2.1":              1.2858, // ￥0.018 / 1k tokens
-	"SparkDesk-v3.1":              1.2858, // ￥0.018 / 1k tokens
-	"SparkDesk-v3.1-128K":         1.2858, // ￥0.018 / 1k tokens
-	"SparkDesk-v3.5":              1.2858, // ￥0.018 / 1k tokens
-	"SparkDesk-v3.5-32K":          1.2858, // ￥0.018 / 1k tokens
-	"SparkDesk-v4.0":              1.2858, // ￥0.018 / 1k tokens
-	"360GPT_S2_V9":                0.8572, // ¥0.012 / 1k tokens
-	"embedding-bert-512-v1":       0.0715, // ¥0.001 / 1k tokens
-	"embedding_s1_v1":             0.0715, // ¥0.001 / 1k tokens
-	"semantic_similarity_s1_v1":   0.0715, // ¥0.001 / 1k tokens
+	"qwen-turbo":                    0.0003 * RMB,
+	"qwen-turbo-latest":             0.0003 * RMB,
+	"qwen-plus":                     0.0008 * RMB,
+	"qwen-plus-latest":              0.0008 * RMB,
+	"qwen-max":                      0.0024 * RMB,
+	"qwen-max-latest":               0.0024 * RMB,
+	"qwen-max-longcontext":          0.0005 * RMB,
+	"qwen-vl-max":                   0.003 * RMB,
+	"qwen-vl-max-latest":            0.003 * RMB,
+	"qwen-vl-plus":                  0.0015 * RMB,
+	"qwen-vl-plus-latest":           0.0015 * RMB,
+	"qwen-vl-ocr":                   0.005 * RMB,
+	"qwen-vl-ocr-latest":            0.005 * RMB,
+	"qwen-audio-turbo":              1.4286,
+	"qwen-math-plus":                0.004 * RMB,
+	"qwen-math-plus-latest":         0.004 * RMB,
+	"qwen-math-turbo":               0.002 * RMB,
+	"qwen-math-turbo-latest":        0.002 * RMB,
+	"qwen-coder-plus":               0.0035 * RMB,
+	"qwen-coder-plus-latest":        0.0035 * RMB,
+	"qwen-coder-turbo":              0.002 * RMB,
+	"qwen-coder-turbo-latest":       0.002 * RMB,
+	"qwen-mt-plus":                  0.015 * RMB,
+	"qwen-mt-turbo":                 0.001 * RMB,
+	"qwq-32b-preview":               0.002 * RMB,
+	"qwen2.5-72b-instruct":          0.004 * RMB,
+	"qwen2.5-32b-instruct":          0.03 * RMB,
+	"qwen2.5-14b-instruct":          0.001 * RMB,
+	"qwen2.5-7b-instruct":           0.0005 * RMB,
+	"qwen2.5-3b-instruct":           0.006 * RMB,
+	"qwen2.5-1.5b-instruct":         0.0003 * RMB,
+	"qwen2.5-0.5b-instruct":         0.0003 * RMB,
+	"qwen2-72b-instruct":            0.004 * RMB,
+	"qwen2-57b-a14b-instruct":       0.0035 * RMB,
+	"qwen2-7b-instruct":             0.001 * RMB,
+	"qwen2-1.5b-instruct":           0.001 * RMB,
+	"qwen2-0.5b-instruct":           0.001 * RMB,
+	"qwen1.5-110b-chat":             0.007 * RMB,
+	"qwen1.5-72b-chat":              0.005 * RMB,
+	"qwen1.5-32b-chat":              0.0035 * RMB,
+	"qwen1.5-14b-chat":              0.002 * RMB,
+	"qwen1.5-7b-chat":               0.001 * RMB,
+	"qwen1.5-1.8b-chat":             0.001 * RMB,
+	"qwen1.5-0.5b-chat":             0.001 * RMB,
+	"qwen-72b-chat":                 0.02 * RMB,
+	"qwen-14b-chat":                 0.008 * RMB,
+	"qwen-7b-chat":                  0.006 * RMB,
+	"qwen-1.8b-chat":                0.006 * RMB,
+	"qwen-1.8b-longcontext-chat":    0.006 * RMB,
+	"qvq-72b-preview":               0.012 * RMB,
+	"qwen2.5-vl-72b-instruct":       0.016 * RMB,
+	"qwen2.5-vl-7b-instruct":        0.002 * RMB,
+	"qwen2.5-vl-3b-instruct":        0.0012 * RMB,
+	"qwen2-vl-7b-instruct":          0.016 * RMB,
+	"qwen2-vl-2b-instruct":          0.002 * RMB,
+	"qwen-vl-v1":                    0.002 * RMB,
+	"qwen-vl-chat-v1":               0.002 * RMB,
+	"qwen2-audio-instruct":          0.002 * RMB,
+	"qwen-audio-chat":               0.002 * RMB,
+	"qwen2.5-math-72b-instruct":     0.004 * RMB,
+	"qwen2.5-math-7b-instruct":      0.001 * RMB,
+	"qwen2.5-math-1.5b-instruct":    0.001 * RMB,
+	"qwen2-math-72b-instruct":       0.004 * RMB,
+	"qwen2-math-7b-instruct":        0.001 * RMB,
+	"qwen2-math-1.5b-instruct":      0.001 * RMB,
+	"qwen2.5-coder-32b-instruct":    0.002 * RMB,
+	"qwen2.5-coder-14b-instruct":    0.002 * RMB,
+	"qwen2.5-coder-7b-instruct":     0.001 * RMB,
+	"qwen2.5-coder-3b-instruct":     0.001 * RMB,
+	"qwen2.5-coder-1.5b-instruct":   0.001 * RMB,
+	"qwen2.5-coder-0.5b-instruct":   0.001 * RMB,
+	"text-embedding-v1":             0.0007 * RMB, // ￥0.0007 / 1k tokens
+	"text-embedding-v3":             0.0007 * RMB,
+	"text-embedding-v2":             0.0007 * RMB,
+	"text-embedding-async-v2":       0.0007 * RMB,
+	"text-embedding-async-v1":       0.0007 * RMB,
+	"ali-stable-diffusion-xl":       8.00,
+	"ali-stable-diffusion-v1.5":     8.00,
+	"wanx-v1":                       8.00,
+	"deepseek-r1":                   0.002 * RMB,
+	"deepseek-v3":                   0.001 * RMB,
+	"deepseek-r1-distill-qwen-1.5b": 0.001 * RMB,
+	"deepseek-r1-distill-qwen-7b":   0.0005 * RMB,
+	"deepseek-r1-distill-qwen-14b":  0.001 * RMB,
+	"deepseek-r1-distill-qwen-32b":  0.002 * RMB,
+	"deepseek-r1-distill-llama-8b":  0.0005 * RMB,
+	"deepseek-r1-distill-llama-70b": 0.004 * RMB,
+	"SparkDesk":                     1.2858, // ￥0.018 / 1k tokens
+	"SparkDesk-v1.1":                1.2858, // ￥0.018 / 1k tokens
+	"SparkDesk-v2.1":                1.2858, // ￥0.018 / 1k tokens
+	"SparkDesk-v3.1":                1.2858, // ￥0.018 / 1k tokens
+	"SparkDesk-v3.1-128K":           1.2858, // ￥0.018 / 1k tokens
+	"SparkDesk-v3.5":                1.2858, // ￥0.018 / 1k tokens
+	"SparkDesk-v3.5-32K":            1.2858, // ￥0.018 / 1k tokens
+	"SparkDesk-v4.0":                1.2858, // ￥0.018 / 1k tokens
+	"360GPT_S2_V9":                  0.8572, // ¥0.012 / 1k tokens
+	"embedding-bert-512-v1":         0.0715, // ¥0.001 / 1k tokens
+	"embedding_s1_v1":               0.0715, // ¥0.001 / 1k tokens
+	"semantic_similarity_s1_v1":     0.0715, // ¥0.001 / 1k tokens
 	// https://cloud.tencent.com/document/product/1729/97731#e0e6be58-60c8-469f-bdeb-6c264ce3b4d0
 	"hunyuan-turbo":             0.015 * RMB,
 	"hunyuan-large":             0.004 * RMB,
@@ -327,6 +343,9 @@ var ModelRatio = map[string]float64{
 	"deepl-ja": 25.0 / 1000 * USD,
 	// https://console.x.ai/
 	"grok-beta": 5.0 / 1000 * USD,
+	// vertex imagen3
+	// https://cloud.google.com/vertex-ai/generative-ai/pricing#imagen-models
+	"imagen-3.0-generate-001": 0.02 * USD,
 	// replicate charges based on the number of generated images
 	// https://replicate.com/pricing
 	"black-forest-labs/flux-1.1-pro":                0.04 * USD,
@@ -371,6 +390,238 @@ var ModelRatio = map[string]float64{
 	"mistralai/mistral-7b-instruct-v0.2":        0.050 * USD,
 	"mistralai/mistral-7b-v0.1":                 0.050 * USD,
 	"mistralai/mixtral-8x7b-instruct-v0.1":      0.300 * USD,
+	//https://openrouter.ai/models
+	"01-ai/yi-large":                                  1.5,
+	"aetherwiing/mn-starcannon-12b":                   0.6,
+	"ai21/jamba-1-5-large":                            4.0,
+	"ai21/jamba-1-5-mini":                             0.2,
+	"ai21/jamba-instruct":                             0.35,
+	"aion-labs/aion-1.0":                              6.0,
+	"aion-labs/aion-1.0-mini":                         1.2,
+	"aion-labs/aion-rp-llama-3.1-8b":                  0.1,
+	"allenai/llama-3.1-tulu-3-405b":                   5.0,
+	"alpindale/goliath-120b":                          4.6875,
+	"alpindale/magnum-72b":                            1.125,
+	"amazon/nova-lite-v1":                             0.12,
+	"amazon/nova-micro-v1":                            0.07,
+	"amazon/nova-pro-v1":                              1.6,
+	"anthracite-org/magnum-v2-72b":                    1.5,
+	"anthracite-org/magnum-v4-72b":                    1.125,
+	"anthropic/claude-2":                              12.0,
+	"anthropic/claude-2.0":                            12.0,
+	"anthropic/claude-2.0:beta":                       12.0,
+	"anthropic/claude-2.1":                            12.0,
+	"anthropic/claude-2.1:beta":                       12.0,
+	"anthropic/claude-2:beta":                         12.0,
+	"anthropic/claude-3-haiku":                        0.625,
+	"anthropic/claude-3-haiku:beta":                   0.625,
+	"anthropic/claude-3-opus":                         37.5,
+	"anthropic/claude-3-opus:beta":                    37.5,
+	"anthropic/claude-3-sonnet":                       7.5,
+	"anthropic/claude-3-sonnet:beta":                  7.5,
+	"anthropic/claude-3.5-haiku":                      2.0,
+	"anthropic/claude-3.5-haiku-20241022":             2.0,
+	"anthropic/claude-3.5-haiku-20241022:beta":        2.0,
+	"anthropic/claude-3.5-haiku:beta":                 2.0,
+	"anthropic/claude-3.5-sonnet":                     7.5,
+	"anthropic/claude-3.5-sonnet-20240620":            7.5,
+	"anthropic/claude-3.5-sonnet-20240620:beta":       7.5,
+	"anthropic/claude-3.5-sonnet:beta":                7.5,
+	"cognitivecomputations/dolphin-mixtral-8x22b":     0.45,
+	"cognitivecomputations/dolphin-mixtral-8x7b":      0.25,
+	"cohere/command":                                  0.95,
+	"cohere/command-r":                                0.7125,
+	"cohere/command-r-03-2024":                        0.7125,
+	"cohere/command-r-08-2024":                        0.285,
+	"cohere/command-r-plus":                           7.125,
+	"cohere/command-r-plus-04-2024":                   7.125,
+	"cohere/command-r-plus-08-2024":                   4.75,
+	"cohere/command-r7b-12-2024":                      0.075,
+	"databricks/dbrx-instruct":                        0.6,
+	"deepseek/deepseek-chat":                          0.445,
+	"deepseek/deepseek-chat-v2.5":                     1.0,
+	"deepseek/deepseek-chat:free":                     0.0,
+	"deepseek/deepseek-r1":                            1.2,
+	"deepseek/deepseek-r1-distill-llama-70b":          0.345,
+	"deepseek/deepseek-r1-distill-llama-70b:free":     0.0,
+	"deepseek/deepseek-r1-distill-llama-8b":           0.02,
+	"deepseek/deepseek-r1-distill-qwen-1.5b":          0.09,
+	"deepseek/deepseek-r1-distill-qwen-14b":           0.075,
+	"deepseek/deepseek-r1-distill-qwen-32b":           0.09,
+	"deepseek/deepseek-r1:free":                       0.0,
+	"eva-unit-01/eva-llama-3.33-70b":                  3.0,
+	"eva-unit-01/eva-qwen-2.5-32b":                    1.7,
+	"eva-unit-01/eva-qwen-2.5-72b":                    3.0,
+	"google/gemini-2.0-flash-001":                     0.2,
+	"google/gemini-2.0-flash-exp:free":                0.0,
+	"google/gemini-2.0-flash-lite-preview-02-05:free": 0.0,
+	"google/gemini-2.0-flash-thinking-exp-1219:free":  0.0,
+	"google/gemini-2.0-flash-thinking-exp:free":       0.0,
+	"google/gemini-2.0-pro-exp-02-05:free":            0.0,
+	"google/gemini-exp-1206:free":                     0.0,
+	"google/gemini-flash-1.5":                         0.15,
+	"google/gemini-flash-1.5-8b":                      0.075,
+	"google/gemini-flash-1.5-8b-exp":                  0.0,
+	"google/gemini-pro":                               0.75,
+	"google/gemini-pro-1.5":                           2.5,
+	"google/gemini-pro-vision":                        0.75,
+	"google/gemma-2-27b-it":                           0.135,
+	"google/gemma-2-9b-it":                            0.03,
+	"google/gemma-2-9b-it:free":                       0.0,
+	"google/gemma-7b-it":                              0.075,
+	"google/learnlm-1.5-pro-experimental:free":        0.0,
+	"google/palm-2-chat-bison":                        1.0,
+	"google/palm-2-chat-bison-32k":                    1.0,
+	"google/palm-2-codechat-bison":                    1.0,
+	"google/palm-2-codechat-bison-32k":                1.0,
+	"gryphe/mythomax-l2-13b":                          0.0325,
+	"gryphe/mythomax-l2-13b:free":                     0.0,
+	"huggingfaceh4/zephyr-7b-beta:free":               0.0,
+	"infermatic/mn-inferor-12b":                       0.6,
+	"inflection/inflection-3-pi":                      5.0,
+	"inflection/inflection-3-productivity":            5.0,
+	"jondurbin/airoboros-l2-70b":                      0.25,
+	"liquid/lfm-3b":                                   0.01,
+	"liquid/lfm-40b":                                  0.075,
+	"liquid/lfm-7b":                                   0.005,
+	"mancer/weaver":                                   1.125,
+	"meta-llama/llama-2-13b-chat":                     0.11,
+	"meta-llama/llama-2-70b-chat":                     0.45,
+	"meta-llama/llama-3-70b-instruct":                 0.2,
+	"meta-llama/llama-3-8b-instruct":                  0.03,
+	"meta-llama/llama-3-8b-instruct:free":             0.0,
+	"meta-llama/llama-3.1-405b":                       1.0,
+	"meta-llama/llama-3.1-405b-instruct":              0.4,
+	"meta-llama/llama-3.1-70b-instruct":               0.15,
+	"meta-llama/llama-3.1-8b-instruct":                0.025,
+	"meta-llama/llama-3.2-11b-vision-instruct":        0.0275,
+	"meta-llama/llama-3.2-11b-vision-instruct:free":   0.0,
+	"meta-llama/llama-3.2-1b-instruct":                0.005,
+	"meta-llama/llama-3.2-3b-instruct":                0.0125,
+	"meta-llama/llama-3.2-90b-vision-instruct":        0.8,
+	"meta-llama/llama-3.3-70b-instruct":               0.15,
+	"meta-llama/llama-3.3-70b-instruct:free":          0.0,
+	"meta-llama/llama-guard-2-8b":                     0.1,
+	"microsoft/phi-3-medium-128k-instruct":            0.5,
+	"microsoft/phi-3-medium-128k-instruct:free":       0.0,
+	"microsoft/phi-3-mini-128k-instruct":              0.05,
+	"microsoft/phi-3-mini-128k-instruct:free":         0.0,
+	"microsoft/phi-3.5-mini-128k-instruct":            0.05,
+	"microsoft/phi-4":                                 0.07,
+	"microsoft/wizardlm-2-7b":                         0.035,
+	"microsoft/wizardlm-2-8x22b":                      0.25,
+	"minimax/minimax-01":                              0.55,
+	"mistralai/codestral-2501":                        0.45,
+	"mistralai/codestral-mamba":                       0.125,
+	"mistralai/ministral-3b":                          0.02,
+	"mistralai/ministral-8b":                          0.05,
+	"mistralai/mistral-7b-instruct":                   0.0275,
+	"mistralai/mistral-7b-instruct-v0.1":              0.1,
+	"mistralai/mistral-7b-instruct-v0.3":              0.0275,
+	"mistralai/mistral-7b-instruct:free":              0.0,
+	"mistralai/mistral-large":                         3.0,
+	"mistralai/mistral-large-2407":                    3.0,
+	"mistralai/mistral-large-2411":                    3.0,
+	"mistralai/mistral-medium":                        4.05,
+	"mistralai/mistral-nemo":                          0.04,
+	"mistralai/mistral-nemo:free":                     0.0,
+	"mistralai/mistral-small":                         0.3,
+	"mistralai/mistral-small-24b-instruct-2501":       0.07,
+	"mistralai/mistral-small-24b-instruct-2501:free":  0.0,
+	"mistralai/mistral-tiny":                          0.125,
+	"mistralai/mixtral-8x22b-instruct":                0.45,
+	"mistralai/mixtral-8x7b":                          0.3,
+	"mistralai/mixtral-8x7b-instruct":                 0.12,
+	"mistralai/pixtral-12b":                           0.05,
+	"mistralai/pixtral-large-2411":                    3.0,
+	"neversleep/llama-3-lumimaid-70b":                 2.25,
+	"neversleep/llama-3-lumimaid-8b":                  0.5625,
+	"neversleep/llama-3-lumimaid-8b:extended":         0.5625,
+	"neversleep/llama-3.1-lumimaid-70b":               2.25,
+	"neversleep/llama-3.1-lumimaid-8b":                0.5625,
+	"neversleep/noromaid-20b":                         1.125,
+	"nothingiisreal/mn-celeste-12b":                   0.6,
+	"nousresearch/hermes-2-pro-llama-3-8b":            0.02,
+	"nousresearch/hermes-3-llama-3.1-405b":            0.4,
+	"nousresearch/hermes-3-llama-3.1-70b":             0.15,
+	"nousresearch/nous-hermes-2-mixtral-8x7b-dpo":     0.3,
+	"nousresearch/nous-hermes-llama2-13b":             0.085,
+	"nvidia/llama-3.1-nemotron-70b-instruct":          0.15,
+	"nvidia/llama-3.1-nemotron-70b-instruct:free":     0.0,
+	"openai/chatgpt-4o-latest":                        7.5,
+	"openai/gpt-3.5-turbo":                            0.75,
+	"openai/gpt-3.5-turbo-0125":                       0.75,
+	"openai/gpt-3.5-turbo-0613":                       1.0,
+	"openai/gpt-3.5-turbo-1106":                       1.0,
+	"openai/gpt-3.5-turbo-16k":                        2.0,
+	"openai/gpt-3.5-turbo-instruct":                   1.0,
+	"openai/gpt-4":                                    30.0,
+	"openai/gpt-4-0314":                               30.0,
+	"openai/gpt-4-1106-preview":                       15.0,
+	"openai/gpt-4-32k":                                60.0,
+	"openai/gpt-4-32k-0314":                           60.0,
+	"openai/gpt-4-turbo":                              15.0,
+	"openai/gpt-4-turbo-preview":                      15.0,
+	"openai/gpt-4o":                                   5.0,
+	"openai/gpt-4o-2024-05-13":                        7.5,
+	"openai/gpt-4o-2024-08-06":                        5.0,
+	"openai/gpt-4o-2024-11-20":                        5.0,
+	"openai/gpt-4o-mini":                              0.3,
+	"openai/gpt-4o-mini-2024-07-18":                   0.3,
+	"openai/gpt-4o:extended":                          9.0,
+	"openai/o1":                                       30.0,
+	"openai/o1-mini":                                  2.2,
+	"openai/o1-mini-2024-09-12":                       2.2,
+	"openai/o1-preview":                               30.0,
+	"openai/o1-preview-2024-09-12":                    30.0,
+	"openai/o3-mini":                                  2.2,
+	"openai/o3-mini-high":                             2.2,
+	"openchat/openchat-7b":                            0.0275,
+	"openchat/openchat-7b:free":                       0.0,
+	"openrouter/auto":                                 -500000.0,
+	"perplexity/llama-3.1-sonar-huge-128k-online":     2.5,
+	"perplexity/llama-3.1-sonar-large-128k-chat":      0.5,
+	"perplexity/llama-3.1-sonar-large-128k-online":    0.5,
+	"perplexity/llama-3.1-sonar-small-128k-chat":      0.1,
+	"perplexity/llama-3.1-sonar-small-128k-online":    0.1,
+	"perplexity/sonar":                                0.5,
+	"perplexity/sonar-reasoning":                      2.5,
+	"pygmalionai/mythalion-13b":                       0.6,
+	"qwen/qvq-72b-preview":                            0.25,
+	"qwen/qwen-2-72b-instruct":                        0.45,
+	"qwen/qwen-2-7b-instruct":                         0.027,
+	"qwen/qwen-2-7b-instruct:free":                    0.0,
+	"qwen/qwen-2-vl-72b-instruct":                     0.2,
+	"qwen/qwen-2-vl-7b-instruct":                      0.05,
+	"qwen/qwen-2.5-72b-instruct":                      0.2,
+	"qwen/qwen-2.5-7b-instruct":                       0.025,
+	"qwen/qwen-2.5-coder-32b-instruct":                0.08,
+	"qwen/qwen-max":                                   3.2,
+	"qwen/qwen-plus":                                  0.6,
+	"qwen/qwen-turbo":                                 0.1,
+	"qwen/qwen-vl-plus:free":                          0.0,
+	"qwen/qwen2.5-vl-72b-instruct:free":               0.0,
+	"qwen/qwq-32b-preview":                            0.09,
+	"raifle/sorcererlm-8x22b":                         2.25,
+	"sao10k/fimbulvetr-11b-v2":                        0.6,
+	"sao10k/l3-euryale-70b":                           0.4,
+	"sao10k/l3-lunaris-8b":                            0.03,
+	"sao10k/l3.1-70b-hanami-x1":                       1.5,
+	"sao10k/l3.1-euryale-70b":                         0.4,
+	"sao10k/l3.3-euryale-70b":                         0.4,
+	"sophosympatheia/midnight-rose-70b":               0.4,
+	"sophosympatheia/rogue-rose-103b-v0.2:free":       0.0,
+	"teknium/openhermes-2.5-mistral-7b":               0.085,
+	"thedrummer/rocinante-12b":                        0.25,
+	"thedrummer/unslopnemo-12b":                       0.25,
+	"undi95/remm-slerp-l2-13b":                        0.6,
+	"undi95/toppy-m-7b":                               0.035,
+	"undi95/toppy-m-7b:free":                          0.0,
+	"x-ai/grok-2-1212":                                5.0,
+	"x-ai/grok-2-vision-1212":                         5.0,
+	"x-ai/grok-beta":                                  7.5,
+	"x-ai/grok-vision-beta":                           7.5,
+	"xwin-lm/xwin-lm-70b":                             1.875,
 }

 var CompletionRatio = map[string]float64{
--- a/relay/controller/image.go
+++ b/relay/controller/image.go
@@ -19,7 +19,7 @@ import (
 	"github.com/songquanpeng/one-api/relay/adaptor/openai"
 	billingratio "github.com/songquanpeng/one-api/relay/billing/ratio"
 	"github.com/songquanpeng/one-api/relay/channeltype"
-	"github.com/songquanpeng/one-api/relay/meta"
+	metalib "github.com/songquanpeng/one-api/relay/meta"
 	relaymodel "github.com/songquanpeng/one-api/relay/model"
 )

@@ -66,7 +66,7 @@ func getImageSizeRatio(model string, size string) float64 {
 	return 1
 }

-func validateImageRequest(imageRequest *relaymodel.ImageRequest, _ *meta.Meta) *relaymodel.ErrorWithStatusCode {
+func validateImageRequest(imageRequest *relaymodel.ImageRequest, _ *metalib.Meta) *relaymodel.ErrorWithStatusCode {
 	// check prompt length
 	if imageRequest.Prompt == "" {
 		return openai.ErrorWrapper(errors.New("prompt is required"), "prompt_missing", http.StatusBadRequest)
@@ -105,7 +105,7 @@ func getImageCostRatio(imageRequest *relaymodel.ImageRequest) (float64, error) {

 func RelayImageHelper(c *gin.Context, relayMode int) *relaymodel.ErrorWithStatusCode {
 	ctx := c.Request.Context()
-	meta := meta.GetByContext(c)
+	meta := metalib.GetByContext(c)
 	imageRequest, err := getImageRequest(c, meta.Mode)
 	if err != nil {
 		logger.Errorf(ctx, "getImageRequest failed: %s", err.Error())
@@ -117,6 +117,7 @@ func RelayImageHelper(c *gin.Context, relayMode int) *relaymodel.ErrorWithStatus
 	meta.OriginModelName = imageRequest.Model
 	imageRequest.Model, isModelMapped = getMappedModelName(imageRequest.Model, meta.ModelMapping)
 	meta.ActualModelName = imageRequest.Model
+	metalib.Set2Context(c, meta)

 	// model validation
 	bizErr := validateImageRequest(imageRequest, meta)
@@ -156,8 +157,9 @@ func RelayImageHelper(c *gin.Context, relayMode int) *relaymodel.ErrorWithStatus
 	case channeltype.Zhipu,
 		channeltype.Ali,
 		channeltype.Replicate,
+		channeltype.VertextAI,
 		channeltype.Baidu:
-		finalRequest, err := adaptor.ConvertImageRequest(imageRequest)
+		finalRequest, err := adaptor.ConvertImageRequest(c, imageRequest)
 		if err != nil {
 			return openai.ErrorWrapper(err, "convert_image_request_failed", http.StatusInternalServerError)
 		}
--- a/relay/meta/relay_meta.go
+++ b/relay/meta/relay_meta.go
@@ -38,6 +38,10 @@ type Meta struct {
 }

 func GetByContext(c *gin.Context) *Meta {
+	if v, ok := c.Get(ctxkey.Meta); ok {
+		return v.(*Meta)
+	}
+
 	meta := Meta{
 		Mode:               relaymode.GetByPath(c.Request.URL.Path),
 		ChannelType:        c.GetInt(ctxkey.Channel),
@@ -62,5 +66,11 @@ func GetByContext(c *gin.Context) *Meta {
 		meta.BaseURL = channeltype.ChannelBaseURLs[meta.ChannelType]
 	}
 	meta.APIType = channeltype.ToAPIType(meta.ChannelType)
+
+	Set2Context(c, &meta)
 	return &meta
 }
+
+func Set2Context(c *gin.Context, meta *Meta) {
+	c.Set(ctxkey.Meta, meta)
+}
Author	SHA1	Message	Date
Laisky.Cai	a7a31599d4	Merge `009337ccf3` into `7ac553541b`	2025-02-18 14:31:39 +08:00
longkeyy	7ac553541b	feat: update openrouter models and price 20250213 (#2084 ) Some checks failed CI / Unit tests (push) Has been cancelled Details CI / commit_lint (push) Has been cancelled Details	2025-02-16 18:01:59 +08:00
longkeyy	a5c517c27a	feat: update ali models and price 20250213 (#2086 )	2025-02-16 18:01:24 +08:00
Laisky.Cai	009337ccf3	feat: support vertex imagen3	2025-01-12 04:22:21 +00:00