Nutanix AI distribution

2025-12-18 10:39:48 +00:00 · 2024-11-21 23:45:48 +00:00 · 2024-11-21 23:45:48 +00:00 · cb82b1ee9e
commit cb82b1ee9e
parent f2ac4e2a94
10 changed files with 237 additions and 14 deletions
--- a/llama_stack/providers/remote/inference/nutanix/init.py
+++ b/llama_stack/providers/remote/inference/nutanix/init.py
@ -1,4 +1,4 @@
-# Copyright (c) Meta Platforms, Inc. and affiliates.
+# Copyright (c) Nutanix, Inc. and affiliates.
 # All rights reserved.
 #
 # This source code is licensed under the terms described in the LICENSE file in
--- a/llama_stack/providers/remote/inference/nutanix/config.py
+++ b/llama_stack/providers/remote/inference/nutanix/config.py
@ -1,9 +1,11 @@
-# Copyright (c) Meta Platforms, Inc. and affiliates.
+# Copyright (c) Nutanix, Inc. and affiliates.
 # All rights reserved.
 #
 # This source code is licensed under the terms described in the LICENSE file in
 # the root directory of this source tree.

+from typing import Any, Dict, Optional
+
 from llama_models.schema_utils import json_schema_type
 from pydantic import BaseModel, Field

@ -11,10 +13,17 @@ from pydantic import BaseModel, Field
@json_schema_type
 class NutanixImplConfig(BaseModel):
    url: str = Field(
-        default=None,
-        description="The URL of the Nutanix AI endpoint",
+        default="https://ai.nutanix.com/api/v1",
+        description="The URL of the Nutanix AI Endpoint",
    )
-    api_token: str = Field(
+    api_key: Optional[str] = Field(
        default=None,
-        description="The API token of the Nutanix AI endpoint",
+        description="The API key to the Nutanix AI Endpoint",
    )
+
+    @classmethod
+    def sample_run_config(cls) -> Dict[str, Any]:
+        return {
+            "url": "https://ai.nutanix.com/api/v1",
+            "api_key": "${env.NUTANIX_API_KEY}",
+        }
--- a/llama_stack/providers/remote/inference/nutanix/nutanix.py
+++ b/llama_stack/providers/remote/inference/nutanix/nutanix.py
@ -1,4 +1,4 @@
-# Copyright (c) Meta Platforms, Inc. and affiliates.
+# Copyright (c) Nutanix, Inc. and affiliates.
 # All rights reserved.
 #
 # This source code is licensed under the terms described in the LICENSE file in
@ -30,7 +30,7 @@ from llama_stack.providers.utils.inference.prompt_adapter import (
 from .config import NutanixImplConfig


-model_aliases = [
+MODEL_ALIASES = [
    build_model_alias(
        "vllm-llama-3-1",
        CoreModelId.llama3_1_8b_instruct.value,
@ -40,7 +40,7 @@ model_aliases = [

 class NutanixInferenceAdapter(ModelRegistryHelper, Inference):
    def __init__(self, config: NutanixImplConfig) -> None:
-        ModelRegistryHelper.__init__(self, model_aliases)
+        ModelRegistryHelper.__init__(self, MODEL_ALIASES)
        self.config = config
        self.formatter = ChatFormat(Tokenizer.get_instance())

@ -50,6 +50,20 @@ class NutanixInferenceAdapter(ModelRegistryHelper, Inference):
    async def shutdown(self) -> None:
        pass

+    def _get_client(self) -> OpenAI:
+        nutanix_api_key = None
+        if self.config.api_key:
+            nutanix_api_key = self.config.api_key
+        else:
+            provider_data = self.get_request_provider_data()
+            if provider_data is None or not provider_data.nutanix_api_key:
+                raise ValueError(
+                    'Pass Together API Key in the header X-LlamaStack-ProviderData as { "nutanix_api_key": <your api key>}'
+                )
+            nutanix_api_key = provider_data.nutanix_api_key
+
+        return OpenAI(base_url=self.config.url, api_key=nutanix_api_key)
+
    async def completion(
        self,
        model_id: str,
@ -85,7 +99,7 @@ class NutanixInferenceAdapter(ModelRegistryHelper, Inference):
            logprobs=logprobs,
        )

-        client = OpenAI(base_url=self.config.url, api_key=self.config.api_token)
+        client = self._get_client()
        if stream:
            return self._stream_chat_completion(request, client)
        else: