0ef5fcb1c5
Security / Dependency audit (pip-audit) (push) Has been cancelled
Security / CodeQL (javascript-typescript) (push) Has been cancelled
Security / CodeQL (python) (push) Has been cancelled
Security / Secret scan (gitleaks) (push) Has been cancelled
rust / test (ubuntu) (push) Has been cancelled
rust / simulator e2e (macos-latest) (push) Has been cancelled
rust / simulator e2e (ubuntu-latest) (push) Has been cancelled
rust / simulator e2e (windows-latest) (push) Has been cancelled
rust / wheels (aarch64-apple-darwin) (push) Has been cancelled
rust / wheels (x86_64-unknown-linux-gnu) (push) Has been cancelled
rust / wheels (x86_64-apple-darwin) (push) Has been cancelled
rust / audit (push) Has been cancelled
rust / parity (nightly, allowed to fail during Phase 0) (push) Has been cancelled
CI / commitlint (push) Has been skipped
Dev Containers / validate (.devcontainer/devcontainer.json, default) (push) Failing after 0s
Dev Containers / validate (.devcontainer/memory-stack/devcontainer.json, memory-stack) (push) Failing after 0s
Dev Containers / validate-worktree (push) Failing after 0s
CI / changes (push) Failing after 4s
Deploy Documentation / validate (push) Has been skipped
Deploy Documentation / deploy (push) Failing after 1s
Init Native E2E / init-native (ubuntu-latest, claude) (push) Failing after 1s
Init Native E2E / init-native (ubuntu-latest, codex) (push) Failing after 1s
Install Native E2E / install-native (ubuntu-latest) (push) Failing after 1s
OpenCode Plugin / typecheck + build + test (push) Failing after 1s
Init Native E2E / init-native (ubuntu-latest, copilot) (push) Failing after 1s
Release Please / release-please (push) Failing after 1s
Wrap E2E / docker-wrap-e2e (push) Failing after 1s
Wrap Native E2E / wrap-native (ubuntu-latest) (push) Failing after 1s
Init E2E / docker-init-e2e (push) Failing after 4s
Merge Conflicts / merge-conflicts (push) Failing after 4s
CI / lint (push) Has been cancelled
CI / build-wheel (push) Has been cancelled
CI / build-wheel-windows (push) Has been cancelled
CI / prefetch-model (push) Has been cancelled
CI / test-dashboard-ui (push) Has been cancelled
CI / test (1) (push) Has been cancelled
CI / test (2) (push) Has been cancelled
CI / test (3) (push) Has been cancelled
CI / test (4) (push) Has been cancelled
CI / test-extras (push) Has been cancelled
CI / test-agno (push) Has been cancelled
CI / build (push) Has been cancelled
CI / workflow-validation (push) Has been cancelled
CI / docker-native-e2e (push) Has been cancelled
CI / windows-native-wrapper (push) Has been cancelled
CI / macos-native-wrapper (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime-code-nonroot name:code-nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime-code-slim name:code-slim]) (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime-code-slim-nonroot name:code-slim-nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime-nonroot name:nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime-slim name:slim]) (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime-slim-nonroot name:slim-nonroot]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime name:]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime-code name:code]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime-code-nonroot name:code-nonroot]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime-code-slim name:code-slim]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime-code-slim-nonroot name:code-slim-nonroot]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime-nonroot name:nonroot]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime-slim name:slim]) (push) Has been cancelled
Docker / docker-manifest (map[bake_target:runtime-slim-nonroot name:slim-nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime name:]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime-code name:code]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime-code-nonroot name:code-nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime-code-slim name:code-slim]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime-code-slim-nonroot name:code-slim-nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime-nonroot name:nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime-slim name:slim]) (push) Has been cancelled
Docker / docker-build (map[name:amd64 platform:linux/amd64 runs_on:ubuntu-24.04], map[bake_target:runtime-slim-nonroot name:slim-nonroot]) (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime name:]) (push) Has been cancelled
Docker / docker-build (map[name:arm64 platform:linux/arm64 runs_on:ubuntu-24.04-arm], map[bake_target:runtime-code name:code]) (push) Has been cancelled
Docker / promote-latest (push) Has been cancelled
Init Native E2E / init-native (macos-latest, claude) (push) Has been cancelled
Init Native E2E / init-native (macos-latest, codex) (push) Has been cancelled
Init Native E2E / init-native (macos-latest, copilot) (push) Has been cancelled
Install Native E2E / install-native (macos-latest) (push) Has been cancelled
Wrap Native E2E / wrap-native (macos-latest) (push) Has been cancelled
421 lines
14 KiB
Python
421 lines
14 KiB
Python
from __future__ import annotations
|
|
|
|
import logging
|
|
|
|
from headroom.providers.registry import (
|
|
ProviderApiOverrides,
|
|
build_proxy_provider_runtime,
|
|
create_proxy_backend,
|
|
format_backend_status,
|
|
resolve_api_overrides,
|
|
resolve_api_targets,
|
|
)
|
|
from headroom.proxy.models import ProxyConfig
|
|
|
|
|
|
def test_resolve_api_overrides_prefers_explicit_values_over_environment(monkeypatch) -> None:
|
|
monkeypatch.setenv("ANTHROPIC_TARGET_API_URL", "https://env.anthropic.example/v1")
|
|
monkeypatch.setenv("OPENAI_TARGET_API_URL", "https://env.openai.example/v1")
|
|
monkeypatch.setenv("VERTEX_TARGET_API_URL", "https://env-vertex-aiplatform.example/v1")
|
|
|
|
overrides = resolve_api_overrides(
|
|
anthropic_api_url="https://cli.anthropic.example/v1",
|
|
openai_api_url=None,
|
|
gemini_api_url=None,
|
|
cloudcode_api_url=None,
|
|
vertex_api_url="https://cli-vertex-aiplatform.example/v1",
|
|
)
|
|
|
|
assert overrides == ProviderApiOverrides(
|
|
anthropic="https://cli.anthropic.example/v1",
|
|
openai="https://env.openai.example/v1",
|
|
gemini=None,
|
|
cloudcode=None,
|
|
vertex="https://cli-vertex-aiplatform.example/v1",
|
|
)
|
|
|
|
|
|
def test_resolve_api_targets_normalizes_trailing_v1() -> None:
|
|
targets = resolve_api_targets(
|
|
ProviderApiOverrides(
|
|
anthropic="https://anthropic.example/v1/",
|
|
openai="https://openai.example/v1",
|
|
gemini="https://gemini.example/v1",
|
|
cloudcode="https://cloudcode.example/v1/",
|
|
vertex="https://vertex.example/v1/",
|
|
)
|
|
)
|
|
|
|
assert targets.anthropic == "https://anthropic.example"
|
|
assert targets.openai == "https://openai.example"
|
|
assert targets.gemini == "https://gemini.example"
|
|
assert targets.cloudcode == "https://cloudcode.example"
|
|
assert targets.vertex == "https://vertex.example"
|
|
|
|
|
|
def test_proxy_config_exposes_provider_api_overrides() -> None:
|
|
config = ProxyConfig(
|
|
anthropic_api_url="https://anthropic.example",
|
|
openai_api_url="https://openai.example",
|
|
gemini_api_url=None,
|
|
cloudcode_api_url="https://cloudcode.example",
|
|
vertex_api_url="https://vertex.example",
|
|
)
|
|
|
|
assert config.provider_api_overrides == ProviderApiOverrides(
|
|
anthropic="https://anthropic.example",
|
|
openai="https://openai.example",
|
|
gemini=None,
|
|
cloudcode="https://cloudcode.example",
|
|
vertex="https://vertex.example",
|
|
)
|
|
|
|
|
|
def test_format_backend_status_for_anyllm() -> None:
|
|
assert (
|
|
format_backend_status(
|
|
backend="anyllm",
|
|
anyllm_provider="groq",
|
|
bedrock_region="us-central1",
|
|
)
|
|
== "Groq via any-llm"
|
|
)
|
|
|
|
|
|
def test_format_backend_status_for_anthropic_direct() -> None:
|
|
assert (
|
|
format_backend_status(
|
|
backend="anthropic",
|
|
anyllm_provider="ignored",
|
|
bedrock_region=None,
|
|
)
|
|
== "ANTHROPIC (direct API)"
|
|
)
|
|
|
|
|
|
def test_proxy_provider_runtime_routes_model_metadata_and_passthrough() -> None:
|
|
runtime = build_proxy_provider_runtime(ProxyConfig())
|
|
|
|
assert runtime.model_metadata_provider({"x-api-key": "test"}) == "anthropic"
|
|
assert runtime.model_metadata_provider({}) == "openai"
|
|
assert (
|
|
runtime.select_passthrough_base_url({"x-api-key": "test"}) == runtime.api_targets.anthropic
|
|
)
|
|
assert (
|
|
runtime.select_passthrough_base_url({"x-goog-api-key": "test"})
|
|
== runtime.api_targets.gemini
|
|
)
|
|
assert runtime.select_passthrough_base_url({"api-key": "azure", "x-headroom-base-url": ""}) == (
|
|
runtime.api_targets.openai
|
|
)
|
|
|
|
|
|
def test_create_proxy_backend_handles_missing_litellm_backend(caplog) -> None:
|
|
logger = logging.getLogger("test")
|
|
|
|
with caplog.at_level(logging.WARNING):
|
|
missing = create_proxy_backend(
|
|
backend="bedrock",
|
|
anyllm_provider="ignored",
|
|
bedrock_region="us-east-1",
|
|
logger=logger,
|
|
litellm_backend_cls=lambda provider, region, profile_name=None: (_ for _ in ()).throw(
|
|
ImportError("missing")
|
|
),
|
|
)
|
|
|
|
assert missing is None
|
|
assert "LiteLLM backend not available" in caplog.text
|
|
|
|
|
|
def test_create_proxy_backend_logs_structured_failure_details(caplog) -> None:
|
|
logger = logging.getLogger("test")
|
|
|
|
with caplog.at_level(logging.ERROR):
|
|
missing = create_proxy_backend(
|
|
backend="bedrock",
|
|
anyllm_provider="ignored",
|
|
bedrock_region="us-east-1",
|
|
logger=logger,
|
|
litellm_backend_cls=lambda provider, region, profile_name=None: (_ for _ in ()).throw(
|
|
RuntimeError("boom")
|
|
),
|
|
)
|
|
|
|
assert missing is None
|
|
assert "backend initialization failed: backend=litellm-bedrock provider=bedrock error=boom" in (
|
|
caplog.text
|
|
)
|
|
|
|
|
|
def test_proxy_provider_runtime_loaders_cache_backend_types(monkeypatch) -> None:
|
|
import headroom.providers.registry as registry
|
|
|
|
anyllm_loads = 0
|
|
litellm_loads = 0
|
|
|
|
class FakeAnyLLMBackend:
|
|
pass
|
|
|
|
class FakeLiteLLMBackend:
|
|
pass
|
|
|
|
def fake_import(name, globals=None, locals=None, fromlist=(), level=0):
|
|
nonlocal anyllm_loads, litellm_loads
|
|
if name == "headroom.backends.anyllm":
|
|
anyllm_loads += 1
|
|
return type("Module", (), {"AnyLLMBackend": FakeAnyLLMBackend})()
|
|
if name == "headroom.backends.litellm":
|
|
litellm_loads += 1
|
|
return type("Module", (), {"LiteLLMBackend": FakeLiteLLMBackend})()
|
|
raise AssertionError(name)
|
|
|
|
monkeypatch.setattr(registry, "AnyLLMBackendType", None)
|
|
monkeypatch.setattr(registry, "LiteLLMBackendType", None)
|
|
monkeypatch.setattr("builtins.__import__", fake_import)
|
|
|
|
assert registry._load_anyllm_backend() is FakeAnyLLMBackend
|
|
assert registry._load_anyllm_backend() is FakeAnyLLMBackend
|
|
assert registry._load_litellm_backend() is FakeLiteLLMBackend
|
|
assert registry._load_litellm_backend() is FakeLiteLLMBackend
|
|
assert anyllm_loads == 1
|
|
assert litellm_loads == 1
|
|
|
|
|
|
def test_proxy_provider_runtime_transport_helpers_handle_missing_usage() -> None:
|
|
import headroom.providers.registry as registry
|
|
|
|
class Storage:
|
|
def __init__(self) -> None:
|
|
self.saved = []
|
|
|
|
def save(self, metrics) -> None:
|
|
self.saved.append(metrics)
|
|
|
|
client = type(
|
|
"Client",
|
|
(),
|
|
{
|
|
"_storage": Storage(),
|
|
"_original": type(
|
|
"Original",
|
|
(),
|
|
{
|
|
"chat": type(
|
|
"Chat",
|
|
(),
|
|
{
|
|
"completions": type(
|
|
"Completions",
|
|
(),
|
|
{
|
|
"create": staticmethod(
|
|
lambda **kwargs: type("Resp", (), {"usage": None})()
|
|
)
|
|
},
|
|
)()
|
|
},
|
|
)(),
|
|
"messages": type(
|
|
"Messages",
|
|
(),
|
|
{
|
|
"create": staticmethod(
|
|
lambda **kwargs: type("Resp", (), {"usage": None})()
|
|
)
|
|
},
|
|
)(),
|
|
},
|
|
)(),
|
|
},
|
|
)()
|
|
openai_metrics = type("Metrics", (), {"tokens_output": 0, "cached_tokens": 0})()
|
|
anthropic_metrics = type("Metrics", (), {"tokens_output": 0, "cached_tokens": 0})()
|
|
|
|
registry._call_openai_transport(
|
|
client,
|
|
model="gpt-4o",
|
|
messages=[],
|
|
stream=False,
|
|
metrics=openai_metrics,
|
|
)
|
|
registry._call_anthropic_transport(
|
|
client,
|
|
model="claude",
|
|
messages=[],
|
|
stream=False,
|
|
metrics=anthropic_metrics,
|
|
)
|
|
|
|
assert openai_metrics.tokens_output == 0
|
|
assert openai_metrics.cached_tokens == 0
|
|
assert anthropic_metrics.tokens_output == 0
|
|
assert anthropic_metrics.cached_tokens == 0
|
|
assert len(client._storage.saved) == 2
|
|
|
|
|
|
def test_proxy_provider_runtime_transport_helpers_handle_usage_without_optional_cache_fields() -> (
|
|
None
|
|
):
|
|
import headroom.providers.registry as registry
|
|
|
|
class Storage:
|
|
def __init__(self) -> None:
|
|
self.saved = []
|
|
|
|
def save(self, metrics) -> None:
|
|
self.saved.append(metrics)
|
|
|
|
client = type(
|
|
"Client",
|
|
(),
|
|
{
|
|
"_storage": Storage(),
|
|
"_original": type(
|
|
"Original",
|
|
(),
|
|
{
|
|
"chat": type(
|
|
"Chat",
|
|
(),
|
|
{
|
|
"completions": type(
|
|
"Completions",
|
|
(),
|
|
{
|
|
"create": staticmethod(
|
|
lambda **kwargs: type(
|
|
"Resp",
|
|
(),
|
|
{
|
|
"usage": type(
|
|
"Usage",
|
|
(),
|
|
{"completion_tokens": 7},
|
|
)()
|
|
},
|
|
)()
|
|
)
|
|
},
|
|
)()
|
|
},
|
|
)(),
|
|
"messages": type(
|
|
"Messages",
|
|
(),
|
|
{
|
|
"create": staticmethod(
|
|
lambda **kwargs: type(
|
|
"Resp",
|
|
(),
|
|
{
|
|
"usage": type(
|
|
"Usage",
|
|
(),
|
|
{"output_tokens": 5},
|
|
)()
|
|
},
|
|
)()
|
|
)
|
|
},
|
|
)(),
|
|
},
|
|
)(),
|
|
},
|
|
)()
|
|
openai_metrics = type("Metrics", (), {"tokens_output": 0, "cached_tokens": 0})()
|
|
anthropic_metrics = type("Metrics", (), {"tokens_output": 0, "cached_tokens": 0})()
|
|
|
|
registry._call_openai_transport(
|
|
client,
|
|
model="gpt-4o",
|
|
messages=[],
|
|
stream=False,
|
|
metrics=openai_metrics,
|
|
)
|
|
registry._call_anthropic_transport(
|
|
client,
|
|
model="claude",
|
|
messages=[],
|
|
stream=False,
|
|
metrics=anthropic_metrics,
|
|
)
|
|
|
|
assert openai_metrics.tokens_output == 7
|
|
assert openai_metrics.cached_tokens == 0
|
|
assert anthropic_metrics.tokens_output == 5
|
|
assert anthropic_metrics.cached_tokens == 0
|
|
assert len(client._storage.saved) == 2
|
|
|
|
|
|
def test_proxy_provider_runtime_openai_transport_handles_prompt_details_without_cached_tokens() -> (
|
|
None
|
|
):
|
|
import headroom.providers.registry as registry
|
|
|
|
class Storage:
|
|
def __init__(self) -> None:
|
|
self.saved = []
|
|
|
|
def save(self, metrics) -> None:
|
|
self.saved.append(metrics)
|
|
|
|
client = type(
|
|
"Client",
|
|
(),
|
|
{
|
|
"_storage": Storage(),
|
|
"_original": type(
|
|
"Original",
|
|
(),
|
|
{
|
|
"chat": type(
|
|
"Chat",
|
|
(),
|
|
{
|
|
"completions": type(
|
|
"Completions",
|
|
(),
|
|
{
|
|
"create": staticmethod(
|
|
lambda **kwargs: type(
|
|
"Resp",
|
|
(),
|
|
{
|
|
"usage": type(
|
|
"Usage",
|
|
(),
|
|
{
|
|
"completion_tokens": 9,
|
|
"prompt_tokens_details": type(
|
|
"Details",
|
|
(),
|
|
{},
|
|
)(),
|
|
},
|
|
)()
|
|
},
|
|
)()
|
|
)
|
|
},
|
|
)()
|
|
},
|
|
)()
|
|
},
|
|
)(),
|
|
},
|
|
)()
|
|
metrics = type("Metrics", (), {"tokens_output": 0, "cached_tokens": 0})()
|
|
|
|
registry._call_openai_transport(
|
|
client,
|
|
model="gpt-4o",
|
|
messages=[],
|
|
stream=False,
|
|
metrics=metrics,
|
|
)
|
|
|
|
assert metrics.tokens_output == 9
|
|
assert metrics.cached_tokens == 0
|
|
assert len(client._storage.saved) == 1
|