2024-08-11 12:17:54 -07:00
|
|
|
from shared_configs.enums import EmbeddingProvider
|
2024-07-14 23:31:06 -07:00
|
|
|
from shared_configs.enums import EmbedTextType
|
2024-07-14 10:19:53 -07:00
|
|
|
|
|
|
|
|
2024-04-07 21:25:06 -07:00
|
|
|
MODEL_WARM_UP_STRING = "hi " * 512
|
2025-03-13 10:35:45 -07:00
|
|
|
INFORMATION_CONTENT_MODEL_WARM_UP_STRING = "hi " * 16
|
2024-07-14 10:19:53 -07:00
|
|
|
DEFAULT_OPENAI_MODEL = "text-embedding-3-small"
|
|
|
|
DEFAULT_COHERE_MODEL = "embed-english-light-v3.0"
|
|
|
|
DEFAULT_VOYAGE_MODEL = "voyage-large-2-instruct"
|
2025-03-04 17:59:46 -08:00
|
|
|
DEFAULT_VERTEX_MODEL = "text-embedding-005"
|
2024-07-14 10:19:53 -07:00
|
|
|
|
|
|
|
|
|
|
|
class EmbeddingModelTextType:
|
|
|
|
PROVIDER_TEXT_TYPE_MAP = {
|
|
|
|
EmbeddingProvider.COHERE: {
|
|
|
|
EmbedTextType.QUERY: "search_query",
|
|
|
|
EmbedTextType.PASSAGE: "search_document",
|
|
|
|
},
|
|
|
|
EmbeddingProvider.VOYAGE: {
|
|
|
|
EmbedTextType.QUERY: "query",
|
|
|
|
EmbedTextType.PASSAGE: "document",
|
|
|
|
},
|
|
|
|
EmbeddingProvider.GOOGLE: {
|
|
|
|
EmbedTextType.QUERY: "RETRIEVAL_QUERY",
|
|
|
|
EmbedTextType.PASSAGE: "RETRIEVAL_DOCUMENT",
|
|
|
|
},
|
|
|
|
}
|
|
|
|
|
|
|
|
@staticmethod
|
|
|
|
def get_type(provider: EmbeddingProvider, text_type: EmbedTextType) -> str:
|
|
|
|
return EmbeddingModelTextType.PROVIDER_TEXT_TYPE_MAP[provider][text_type]
|
2025-02-07 16:59:02 -08:00
|
|
|
|
|
|
|
|
|
|
|
class GPUStatus:
|
|
|
|
CUDA = "cuda"
|
|
|
|
MAC_MPS = "mps"
|
|
|
|
NONE = "none"
|