mirror of
https://github.com/danswer-ai/danswer.git
synced 2025-03-26 17:51:54 +01:00
31 lines
1018 B
Python
31 lines
1018 B
Python
from shared_configs.enums import EmbeddingProvider
|
|
from shared_configs.enums import EmbedTextType
|
|
|
|
|
|
MODEL_WARM_UP_STRING = "hi " * 512
|
|
DEFAULT_OPENAI_MODEL = "text-embedding-3-small"
|
|
DEFAULT_COHERE_MODEL = "embed-english-light-v3.0"
|
|
DEFAULT_VOYAGE_MODEL = "voyage-large-2-instruct"
|
|
DEFAULT_VERTEX_MODEL = "text-embedding-004"
|
|
|
|
|
|
class EmbeddingModelTextType:
|
|
PROVIDER_TEXT_TYPE_MAP = {
|
|
EmbeddingProvider.COHERE: {
|
|
EmbedTextType.QUERY: "search_query",
|
|
EmbedTextType.PASSAGE: "search_document",
|
|
},
|
|
EmbeddingProvider.VOYAGE: {
|
|
EmbedTextType.QUERY: "query",
|
|
EmbedTextType.PASSAGE: "document",
|
|
},
|
|
EmbeddingProvider.GOOGLE: {
|
|
EmbedTextType.QUERY: "RETRIEVAL_QUERY",
|
|
EmbedTextType.PASSAGE: "RETRIEVAL_DOCUMENT",
|
|
},
|
|
}
|
|
|
|
@staticmethod
|
|
def get_type(provider: EmbeddingProvider, text_type: EmbedTextType) -> str:
|
|
return EmbeddingModelTextType.PROVIDER_TEXT_TYPE_MAP[provider][text_type]
|