426e9eeabd
Voice Workbench / headless workbench (mocked backends) (push) Has been cancelled
Voice Workbench / real acoustic lane (nightly, provisioned only) (push) Has been cancelled
ci / test (push) Has been cancelled
ci / lint-and-format (push) Has been cancelled
ci / build (push) Has been cancelled
ci / dev-startup (push) Has been cancelled
gitleaks / gitleaks (push) Has been cancelled
Markdown Links / Relative Markdown Links (push) Has been cancelled
Quality (Extended) / Homepage Build (PR smoke) (push) Has been cancelled
Quality (Extended) / Comment-only diff guard (push) Has been cancelled
Quality (Extended) / Format + Type Safety Ratchet (push) Has been cancelled
Quality (Extended) / Develop Gate (secret scan + UI determinism) (push) Has been cancelled
Quality (Extended) / Develop Gate (lint) (push) Has been cancelled
Chat shell gestures / Chat shell gesture + parity e2e (push) Has been cancelled
Cloud Gateway Discord / Test (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx @biomejs/biome check packages/lifeops-bench/src, benchmark-lint) (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx vitest run --config packages/lifeops-bench/vitest.config.ts --root packages/lifeops-bench --passWithNoTests, benchmark-tests) (push) Has been cancelled
Build Agent Image / build-and-push (push) Has been cancelled
Dev Smoke / bun run dev onboarding chat (push) Has been cancelled
Dev Smoke / Vite HMR dependency-level smoke (push) Has been cancelled
Electrobun Submodule Guard / electrobun gitlink is fetchable (push) Has been cancelled
Publish @elizaos/example-code / check_npm (push) Has been cancelled
Publish @elizaos/example-code / publish_npm (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / verify_version (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / publish_npm (push) Has been cancelled
Sandbox Live Smoke / Sandbox live smoke (push) Has been cancelled
Snap Build & Test / Build Snap (amd64) (push) Has been cancelled
Snap Build & Test / Build Snap (arm64) (push) Has been cancelled
Test Packaging / elizaos CLI global-install smoke (node + bun) (push) Has been cancelled
Cloud Gateway Webhook / Test (push) Has been cancelled
Cloud Tests / lint-and-types (push) Has been cancelled
Cloud Tests / unit-tests (push) Has been cancelled
Cloud Tests / integration-tests (push) Has been cancelled
Cloud Tests / e2e-tests (push) Has been cancelled
CodeQL Advanced / Analyze (javascript-typescript) (push) Has been cancelled
Deploy Apps Worker (Product 2) / Determine environment (push) Has been cancelled
Deploy Apps Worker (Product 2) / Deploy apps worker to apps-control host (${{ needs.determine-env.outputs.environment }}) (push) Has been cancelled
Deploy Eliza Provisioning Worker / Determine environment (push) Has been cancelled
Deploy Eliza Provisioning Worker / Deploy worker to Hetzner host (${{ needs.determine-env.outputs.environment }} @ ${{ needs.determine-env.outputs.deployment_sha }}) (push) Has been cancelled
Dev Smoke / Classify changed paths (push) Has been cancelled
supply-chain / sbom (push) Has been cancelled
supply-chain / vulnerability-scan (push) Has been cancelled
Build, Push & Deploy to Phala Cloud / build-and-push (push) Has been cancelled
Test Packaging / Validate Packaging Configs (push) Has been cancelled
Test Packaging / Build & Test PyPI Package (push) Has been cancelled
Test Packaging / PyPI on Python ${{ matrix.python }} (push) Has been cancelled
Test Packaging / Pack & Test JS Tarballs (push) Has been cancelled
UI Fixture E2E / ui-fixture-e2e (push) Has been cancelled
UI Fixture E2E / fixture-e2e (push) Has been cancelled
UI Story Gate / story-gate (push) Has been cancelled
vault-ci / test (macos-latest) (push) Has been cancelled
vault-ci / test (ubuntu-latest) (push) Has been cancelled
vault-ci / test (windows-latest) (push) Has been cancelled
vault-ci / app-core wiring tests (push) Has been cancelled
verify-patches / verify patches/CHECKSUMS.sha256 (push) Has been cancelled
Voice Benchmark Smoke / voice-emotion fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voiceagentbench fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench-quality unit smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench TypeScript unit (no audio) (push) Has been cancelled
Voice Benchmark Smoke / voice bench smoke summary (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/app-core test bun run --cwd packages/elizaos test bun run --cwd packages/cloud/shared test], app-and-cli) (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/scenario-runner test bun run --cwd packages/vault test bun run --cwd packages/security test bun run --cwd plugins/plugin-coding-tools test], framework-packages) (push) Has been cancelled
Windows CI / windows ([bun run --cwd plugins/plugin-elizacloud test bun run --cwd plugins/plugin-discord test bun run --cwd plugins/plugin-anthropic test bun run --cwd plugins/plugin-openai test bun run --cwd plugins/plugin-app-control test bun run --cwd plugins/pl… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run build --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/agent --concurrency=4 node packages/scripts/run-bash-linux-only.mjs scripts/verify-riscv64-buildpaths.sh node packages/scripts/run… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run typecheck --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/cloud-shared --concurrency=4 bun run --cwd packages/core test bun run --cwd packages/shared test], core-runtime, 75) (push) Has been cancelled
103 lines
3.5 KiB
TypeScript
103 lines
3.5 KiB
TypeScript
/**
|
|
* `emitModelUsageEvent` fires `EventType.MODEL_USED` after each successful
|
|
* Anthropic call, normalizing the SDK's usage shape (prompt/completion vs
|
|
* input/output token names, plus cache read/write counts) into the runtime's
|
|
* usage payload so billing and telemetry consumers see one consistent record.
|
|
* It also emits a structured prompt-cache log line (read/write token counts +
|
|
* hit/miss classification) so operators can diagnose cache warm-up issues
|
|
* without wiring up a MODEL_USED consumer.
|
|
*/
|
|
import type { EventPayload, IAgentRuntime, ModelTypeName } from "@elizaos/core";
|
|
import { EventType, logger } from "@elizaos/core";
|
|
|
|
type ModelUsage = {
|
|
promptTokens?: number;
|
|
completionTokens?: number;
|
|
inputTokens?: number;
|
|
outputTokens?: number;
|
|
totalTokens?: number;
|
|
cacheReadInputTokens?: number;
|
|
cacheCreationInputTokens?: number;
|
|
};
|
|
|
|
export type NormalizedModelUsage = {
|
|
promptTokens: number;
|
|
completionTokens: number;
|
|
totalTokens: number;
|
|
};
|
|
|
|
/**
|
|
* Classify a call's prompt-cache outcome for the structured log line.
|
|
* - "hit": some prompt tokens were served from cache (cacheRead > 0).
|
|
* - "write": nothing was read but a new cache entry was written — the cold
|
|
* first call of a prefix, or a prefix change (the miss an operator hunting
|
|
* cache regressions cares about).
|
|
* - "none": the provider reported no cache activity (prefix below the
|
|
* model's cacheable minimum, caching disabled, or usage shape without
|
|
* cache fields).
|
|
*/
|
|
export function classifyPromptCacheUsage(
|
|
cacheRead: number | undefined,
|
|
cacheWrite: number | undefined
|
|
): "hit" | "write" | "none" {
|
|
if (typeof cacheRead === "number" && cacheRead > 0) {
|
|
return "hit";
|
|
}
|
|
if (typeof cacheWrite === "number" && cacheWrite > 0) {
|
|
return "write";
|
|
}
|
|
return "none";
|
|
}
|
|
|
|
export function emitModelUsageEvent(
|
|
runtime: IAgentRuntime,
|
|
type: ModelTypeName,
|
|
_prompt: string,
|
|
usage: ModelUsage,
|
|
modelName?: string
|
|
): NormalizedModelUsage {
|
|
const promptTokens = usage.promptTokens ?? usage.inputTokens ?? 0;
|
|
const completionTokens = usage.completionTokens ?? usage.outputTokens ?? 0;
|
|
const totalTokens = usage.totalTokens ?? promptTokens + completionTokens;
|
|
const cacheRead = usage.cacheReadInputTokens;
|
|
const cacheWrite = usage.cacheCreationInputTokens;
|
|
const model = modelName?.trim() || String(type);
|
|
|
|
// Structured prompt-cache visibility (#15742): surface cache read/write
|
|
// counts on every call so a cold prefix ("write" with zero reads on a
|
|
// request that was expected to hit) is diagnosable straight from the logs.
|
|
const cacheOutcome = classifyPromptCacheUsage(cacheRead, cacheWrite);
|
|
logger.debug(
|
|
{
|
|
provider: "anthropic",
|
|
model,
|
|
modelType: String(type),
|
|
promptTokens,
|
|
completionTokens,
|
|
cacheReadInputTokens: cacheRead ?? 0,
|
|
cacheCreationInputTokens: cacheWrite ?? 0,
|
|
cacheOutcome,
|
|
},
|
|
`[Anthropic] prompt cache ${cacheOutcome}: read=${cacheRead ?? 0} write=${cacheWrite ?? 0} prompt=${promptTokens} (${model})`
|
|
);
|
|
|
|
runtime.emitEvent(EventType.MODEL_USED, {
|
|
runtime,
|
|
source: "anthropic",
|
|
provider: "anthropic",
|
|
type,
|
|
model,
|
|
modelName: model,
|
|
modelLabel: String(type),
|
|
tokens: {
|
|
prompt: promptTokens,
|
|
completion: completionTokens,
|
|
total: totalTokens,
|
|
...(cacheRead !== undefined ? { cacheRead } : {}),
|
|
...(cacheWrite !== undefined ? { cacheWrite } : {}),
|
|
},
|
|
} as EventPayload);
|
|
|
|
return { promptTokens, completionTokens, totalTokens };
|
|
}
|