fix: thread extra_headers through the call_llm split

The PR's concurrency wrapper splits call_llm into a semaphore-guarded
entry + _call_llm_impl; main added extra_headers to call_llm's
signature after the PR's base, so the split has to forward it too
(dropped silently otherwise — Azure Foundry and custom-endpoint
callers set it).
This commit is contained in:
kshitijk4poor
2026-08-03 22:53:54 +05:30
committed by kshitij
parent 23f8ae32c0
commit ddae511ab1

View File

@@ -8546,6 +8546,7 @@ def call_llm(
timeout=timeout,
extra_body=extra_body,
reasoning_config=reasoning_config,
extra_headers=extra_headers,
api_mode=api_mode,
stream=stream,
stream_options=stream_options,
@@ -8590,6 +8591,7 @@ def _call_llm_impl(
timeout: float = None,
extra_body: dict = None,
reasoning_config: Optional[dict] = None,
extra_headers: Optional[Dict[str, str]] = None,
api_mode: str = None,
stream: bool = False,
stream_options: dict = None,