Merge branch 'main' into groq

2025-12-17 07:12:36 +00:00 · 2024-11-26 12:28:31 -06:00 · 2024-11-26 12:28:31 -06:00 · bc427b3081
commit bc427b3081
parent 7d7d1e6ea1 d3956a1d22
9 changed files with 43 additions and 8 deletions
--- a/llama_stack/providers/tests/inference/fixtures.py
+++ b/llama_stack/providers/tests/inference/fixtures.py
@ -20,6 +20,7 @@ from llama_stack.providers.remote.inference.bedrock import BedrockConfig
 from llama_stack.providers.remote.inference.fireworks import FireworksImplConfig
 from llama_stack.providers.remote.inference.nvidia import NVIDIAConfig
 from llama_stack.providers.remote.inference.ollama import OllamaImplConfig
+from llama_stack.providers.remote.inference.tgi import TGIImplConfig
 from llama_stack.providers.remote.inference.together import TogetherImplConfig
 from llama_stack.providers.remote.inference.vllm import VLLMInferenceAdapterConfig
 from llama_stack.providers.remote.inference.groq import GroqImplConfig
@ -172,6 +173,21 @@ def inference_groq() -> ProviderFixture:
        ),
    )

+@pytest.fixture(scope="session")
+def inference_tgi() -> ProviderFixture:
+    return ProviderFixture(
+        providers=[
+            Provider(
+                provider_id="tgi",
+                provider_type="remote::tgi",
+                config=TGIImplConfig(
+                    url=get_env_or_fail("TGI_URL"),
+                    api_token=os.getenv("TGI_API_TOKEN", None),
+                ).model_dump(),
+            )
+        ],
+    )
+

 def get_model_short_name(model_name: str) -> str:
    """Convert model name to a short test identifier.
@ -208,6 +224,7 @@ INFERENCE_FIXTURES = [
    "bedrock",
    "nvidia",
    "groq",
+    "tgi",
 ]