test_huggingface_hub_model.py 4.3 KB

123456789101112131415161718192021222324252627282930313233343536373839404142434445464748495051525354555657585960616263646566676869707172737475767778798081828384858687888990919293949596979899100101102103104105106107108109110111112113114115116117118119120121122123124125126127128
  1. import json
  2. import os
  3. from unittest.mock import patch, MagicMock
  4. from langchain.schema import Generation
  5. from core.model_providers.models.entity.message import PromptMessage, MessageType
  6. from core.model_providers.models.entity.model_params import ModelKwargs, ModelType
  7. from core.model_providers.models.llm.huggingface_hub_model import HuggingfaceHubModel
  8. from core.model_providers.providers.huggingface_hub_provider import HuggingfaceHubProvider
  9. from models.provider import Provider, ProviderType, ProviderModel
  10. def get_mock_provider():
  11. return Provider(
  12. id='provider_id',
  13. tenant_id='tenant_id',
  14. provider_name='huggingface_hub',
  15. provider_type=ProviderType.CUSTOM.value,
  16. encrypted_config='',
  17. is_valid=True,
  18. )
  19. def get_mock_model(model_name, huggingfacehub_api_type, mocker):
  20. model_kwargs = ModelKwargs(
  21. max_tokens=10,
  22. temperature=0.01
  23. )
  24. valid_api_key = os.environ['HUGGINGFACE_API_KEY']
  25. endpoint_url = os.environ['HUGGINGFACE_ENDPOINT_URL']
  26. model_provider = HuggingfaceHubProvider(provider=get_mock_provider())
  27. credentials = {
  28. 'huggingfacehub_api_type': huggingfacehub_api_type,
  29. 'huggingfacehub_api_token': valid_api_key
  30. }
  31. if huggingfacehub_api_type == 'inference_endpoints':
  32. credentials['huggingfacehub_endpoint_url'] = endpoint_url
  33. mock_query = MagicMock()
  34. mock_query.filter.return_value.first.return_value = ProviderModel(
  35. provider_name='huggingface_hub',
  36. model_name=model_name,
  37. model_type=ModelType.TEXT_GENERATION.value,
  38. encrypted_config=json.dumps(credentials),
  39. is_valid=True,
  40. )
  41. mocker.patch('extensions.ext_database.db.session.query', return_value=mock_query)
  42. return HuggingfaceHubModel(
  43. model_provider=model_provider,
  44. name=model_name,
  45. model_kwargs=model_kwargs
  46. )
  47. def decrypt_side_effect(tenant_id, encrypted_api_key):
  48. return encrypted_api_key
  49. @patch('huggingface_hub.hf_api.ModelInfo')
  50. @patch('core.helper.encrypter.decrypt_token', side_effect=decrypt_side_effect)
  51. def test_hosted_inference_api_get_num_tokens(mock_decrypt, mock_model_info, mocker):
  52. mock_model_info.return_value = MagicMock(pipeline_tag='text2text-generation')
  53. mocker.patch('langchain.llms.huggingface_hub.HuggingFaceHub._call', return_value="abc")
  54. model = get_mock_model(
  55. 'tiiuae/falcon-40b',
  56. 'hosted_inference_api',
  57. mocker
  58. )
  59. rst = model.get_num_tokens([
  60. PromptMessage(type=MessageType.HUMAN, content='Who is your manufacturer?')
  61. ])
  62. assert rst == 5
  63. @patch('huggingface_hub.hf_api.ModelInfo')
  64. @patch('core.helper.encrypter.decrypt_token', side_effect=decrypt_side_effect)
  65. def test_inference_endpoints_get_num_tokens(mock_decrypt, mock_model_info, mocker):
  66. mock_model_info.return_value = MagicMock(pipeline_tag='text2text-generation')
  67. mocker.patch('langchain.llms.huggingface_hub.HuggingFaceHub._call', return_value="abc")
  68. model = get_mock_model(
  69. '',
  70. 'inference_endpoints',
  71. mocker
  72. )
  73. rst = model.get_num_tokens([
  74. PromptMessage(type=MessageType.HUMAN, content='Who is your manufacturer?')
  75. ])
  76. assert rst == 5
  77. @patch('core.helper.encrypter.decrypt_token', side_effect=decrypt_side_effect)
  78. def test_hosted_inference_api_run(mock_decrypt, mocker):
  79. mocker.patch('core.model_providers.providers.base.BaseModelProvider.update_last_used', return_value=None)
  80. model = get_mock_model(
  81. 'google/flan-t5-base',
  82. 'hosted_inference_api',
  83. mocker
  84. )
  85. rst = model.run(
  86. [PromptMessage(content='Human: Are you Really Human? you MUST only answer `y` or `n`? \nAssistant: ')],
  87. stop=['\nHuman:'],
  88. )
  89. assert len(rst.content) > 0
  90. assert rst.content.strip() == 'n'
  91. @patch('core.helper.encrypter.decrypt_token', side_effect=decrypt_side_effect)
  92. def test_inference_endpoints_run(mock_decrypt, mocker):
  93. mocker.patch('core.model_providers.providers.base.BaseModelProvider.update_last_used', return_value=None)
  94. model = get_mock_model(
  95. '',
  96. 'inference_endpoints',
  97. mocker
  98. )
  99. rst = model.run(
  100. [PromptMessage(content='Answer the following yes/no question. Can you write a whole Haiku in a single tweet?')],
  101. )
  102. assert len(rst.content) > 0