More idiomatic REST API (#765)

# What does this PR do? This PR changes our API to follow more idiomatic REST API approaches of having paths being resources and methods indicating the action being performed. Changes made to generator: 1) removed the prefix check of "get" as its not required and is actually needed for other method types too 2) removed _ check on path since variables can have "_" ## Test Plan LLAMA_STACK_BASE_URL=http://localhost:5000 pytest -v tests/client-sdk/agents/test_agents.py
2025-01-15 13:20:09 -08:00 · 2025-01-15 13:20:09 -08:00 · 7fb2c1c48d
commit 7fb2c1c48d
parent 6deef1ece0
29 changed files with 2144 additions and 1917 deletions
--- a/llama_stack/apis/inference/inference.py
+++ b/llama_stack/apis/inference/inference.py
@ -291,7 +291,7 @@ class ModelStore(Protocol):
 class Inference(Protocol):
    model_store: ModelStore

-    @webmethod(route="/inference/completion")
+    @webmethod(route="/inference/completion", method="POST")
    async def completion(
        self,
        model_id: str,
@ -302,7 +302,7 @@ class Inference(Protocol):
        logprobs: Optional[LogProbConfig] = None,
    ) -> Union[CompletionResponse, AsyncIterator[CompletionResponseStreamChunk]]: ...

-    @webmethod(route="/inference/chat-completion")
+    @webmethod(route="/inference/chat-completion", method="POST")
    async def chat_completion(
        self,
        model_id: str,
@ -319,7 +319,7 @@ class Inference(Protocol):
        ChatCompletionResponse, AsyncIterator[ChatCompletionResponseStreamChunk]
    ]: ...

-    @webmethod(route="/inference/embeddings")
+    @webmethod(route="/inference/embeddings", method="POST")
    async def embeddings(
        self,
        model_id: str,