test_xinference_embedding.py 2.1 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566
  1. import json
  2. import os
  3. from unittest.mock import patch, MagicMock
  4. from core.model_providers.models.embedding.xinference_embedding import XinferenceEmbedding
  5. from core.model_providers.models.entity.model_params import ModelType
  6. from core.model_providers.providers.xinference_provider import XinferenceProvider
  7. from models.provider import Provider, ProviderType, ProviderModel
  8. def get_mock_provider():
  9. return Provider(
  10. id='provider_id',
  11. tenant_id='tenant_id',
  12. provider_name='xinference',
  13. provider_type=ProviderType.CUSTOM.value,
  14. encrypted_config='',
  15. is_valid=True,
  16. )
  17. def get_mock_embedding_model(mocker):
  18. model_name = 'vicuna-v1.3'
  19. server_url = os.environ['XINFERENCE_SERVER_URL']
  20. model_uid = os.environ['XINFERENCE_MODEL_UID']
  21. model_provider = XinferenceProvider(provider=get_mock_provider())
  22. mock_query = MagicMock()
  23. mock_query.filter.return_value.first.return_value = ProviderModel(
  24. provider_name='xinference',
  25. model_name=model_name,
  26. model_type=ModelType.EMBEDDINGS.value,
  27. encrypted_config=json.dumps({
  28. 'server_url': server_url,
  29. 'model_uid': model_uid
  30. }),
  31. is_valid=True,
  32. )
  33. mocker.patch('extensions.ext_database.db.session.query', return_value=mock_query)
  34. return XinferenceEmbedding(
  35. model_provider=model_provider,
  36. name=model_name
  37. )
  38. def decrypt_side_effect(tenant_id, encrypted_api_key):
  39. return encrypted_api_key
  40. @patch('core.helper.encrypter.decrypt_token', side_effect=decrypt_side_effect)
  41. def test_embed_documents(mock_decrypt, mocker):
  42. embedding_model = get_mock_embedding_model(mocker)
  43. rst = embedding_model.client.embed_documents(['test', 'test1'])
  44. assert isinstance(rst, list)
  45. assert len(rst) == 2
  46. assert len(rst[0]) == 4096
  47. @patch('core.helper.encrypter.decrypt_token', side_effect=decrypt_side_effect)
  48. def test_embed_query(mock_decrypt, mocker):
  49. embedding_model = get_mock_embedding_model(mocker)
  50. rst = embedding_model.client.embed_query('test')
  51. assert isinstance(rst, list)
  52. assert len(rst) == 4096