fix: don't pass default response format in Responses

# What does this PR do? ## Test Plan
2025-12-05 02:17:31 +00:00 · 2025-09-30 11:24:27 -07:00 · 2025-09-30 11:24:27 -07:00 · 0cc072dcaf
commit 0cc072dcaf
parent 6cce553c93
1 changed files with 4 additions and 1 deletions
--- a/llama_stack/providers/inline/agents/meta_reference/responses/streaming.py
+++ b/llama_stack/providers/inline/agents/meta_reference/responses/streaming.py
@ -127,13 +127,16 @@ class StreamingResponseOrchestrator:
        messages = self.ctx.messages.copy()

        while True:
+            # Text is the default response format for chat completion so don't need to pass it
+            # (some providers don't support non-empty response_format when tools are present)
+            response_format = None if self.ctx.response_format.type == "text" else self.ctx.response_format
            completion_result = await self.inference_api.openai_chat_completion(
                model=self.ctx.model,
                messages=messages,
                tools=self.ctx.chat_tools,
                stream=True,
                temperature=self.ctx.temperature,
-                response_format=self.ctx.response_format,
+                response_format=response_format,
            )

            # Process streaming chunks and build complete response