426e9eeabd
Voice Workbench / headless workbench (mocked backends) (push) Has been cancelled
Voice Workbench / real acoustic lane (nightly, provisioned only) (push) Has been cancelled
ci / test (push) Has been cancelled
ci / lint-and-format (push) Has been cancelled
ci / build (push) Has been cancelled
ci / dev-startup (push) Has been cancelled
gitleaks / gitleaks (push) Has been cancelled
Markdown Links / Relative Markdown Links (push) Has been cancelled
Quality (Extended) / Homepage Build (PR smoke) (push) Has been cancelled
Quality (Extended) / Comment-only diff guard (push) Has been cancelled
Quality (Extended) / Format + Type Safety Ratchet (push) Has been cancelled
Quality (Extended) / Develop Gate (secret scan + UI determinism) (push) Has been cancelled
Quality (Extended) / Develop Gate (lint) (push) Has been cancelled
Chat shell gestures / Chat shell gesture + parity e2e (push) Has been cancelled
Cloud Gateway Discord / Test (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx @biomejs/biome check packages/lifeops-bench/src, benchmark-lint) (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx vitest run --config packages/lifeops-bench/vitest.config.ts --root packages/lifeops-bench --passWithNoTests, benchmark-tests) (push) Has been cancelled
Build Agent Image / build-and-push (push) Has been cancelled
Dev Smoke / bun run dev onboarding chat (push) Has been cancelled
Dev Smoke / Vite HMR dependency-level smoke (push) Has been cancelled
Electrobun Submodule Guard / electrobun gitlink is fetchable (push) Has been cancelled
Publish @elizaos/example-code / check_npm (push) Has been cancelled
Publish @elizaos/example-code / publish_npm (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / verify_version (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / publish_npm (push) Has been cancelled
Sandbox Live Smoke / Sandbox live smoke (push) Has been cancelled
Snap Build & Test / Build Snap (amd64) (push) Has been cancelled
Snap Build & Test / Build Snap (arm64) (push) Has been cancelled
Test Packaging / elizaos CLI global-install smoke (node + bun) (push) Has been cancelled
Cloud Gateway Webhook / Test (push) Has been cancelled
Cloud Tests / lint-and-types (push) Has been cancelled
Cloud Tests / unit-tests (push) Has been cancelled
Cloud Tests / integration-tests (push) Has been cancelled
Cloud Tests / e2e-tests (push) Has been cancelled
CodeQL Advanced / Analyze (javascript-typescript) (push) Has been cancelled
Deploy Apps Worker (Product 2) / Determine environment (push) Has been cancelled
Deploy Apps Worker (Product 2) / Deploy apps worker to apps-control host (${{ needs.determine-env.outputs.environment }}) (push) Has been cancelled
Deploy Eliza Provisioning Worker / Determine environment (push) Has been cancelled
Deploy Eliza Provisioning Worker / Deploy worker to Hetzner host (${{ needs.determine-env.outputs.environment }} @ ${{ needs.determine-env.outputs.deployment_sha }}) (push) Has been cancelled
Dev Smoke / Classify changed paths (push) Has been cancelled
supply-chain / sbom (push) Has been cancelled
supply-chain / vulnerability-scan (push) Has been cancelled
Build, Push & Deploy to Phala Cloud / build-and-push (push) Has been cancelled
Test Packaging / Validate Packaging Configs (push) Has been cancelled
Test Packaging / Build & Test PyPI Package (push) Has been cancelled
Test Packaging / PyPI on Python ${{ matrix.python }} (push) Has been cancelled
Test Packaging / Pack & Test JS Tarballs (push) Has been cancelled
UI Fixture E2E / ui-fixture-e2e (push) Has been cancelled
UI Fixture E2E / fixture-e2e (push) Has been cancelled
UI Story Gate / story-gate (push) Has been cancelled
vault-ci / test (macos-latest) (push) Has been cancelled
vault-ci / test (ubuntu-latest) (push) Has been cancelled
vault-ci / test (windows-latest) (push) Has been cancelled
vault-ci / app-core wiring tests (push) Has been cancelled
verify-patches / verify patches/CHECKSUMS.sha256 (push) Has been cancelled
Voice Benchmark Smoke / voice-emotion fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voiceagentbench fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench-quality unit smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench TypeScript unit (no audio) (push) Has been cancelled
Voice Benchmark Smoke / voice bench smoke summary (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/app-core test bun run --cwd packages/elizaos test bun run --cwd packages/cloud/shared test], app-and-cli) (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/scenario-runner test bun run --cwd packages/vault test bun run --cwd packages/security test bun run --cwd plugins/plugin-coding-tools test], framework-packages) (push) Has been cancelled
Windows CI / windows ([bun run --cwd plugins/plugin-elizacloud test bun run --cwd plugins/plugin-discord test bun run --cwd plugins/plugin-anthropic test bun run --cwd plugins/plugin-openai test bun run --cwd plugins/plugin-app-control test bun run --cwd plugins/pl… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run build --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/agent --concurrency=4 node packages/scripts/run-bash-linux-only.mjs scripts/verify-riscv64-buildpaths.sh node packages/scripts/run… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run typecheck --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/cloud-shared --concurrency=4 bun run --cwd packages/core test bun run --cwd packages/shared test], core-runtime, 75) (push) Has been cancelled
130 lines
4.9 KiB
TypeScript
130 lines
4.9 KiB
TypeScript
/**
|
|
* Live test asserting that the text handlers record their LLM call into the
|
|
* active trajectory context. Post-merge lane, real model.
|
|
*/
|
|
import type { IAgentRuntime } from "@elizaos/core";
|
|
import { runWithTrajectoryContext } from "@elizaos/core";
|
|
import { describe, expect, it } from "vitest";
|
|
|
|
import { describeLive } from "../../../packages/app-core/test/helpers/live-agent-test";
|
|
import { handleTextLarge, handleTextSmall } from "../models/text";
|
|
|
|
interface CapturedLlmCall {
|
|
stepId: string;
|
|
actionType: string;
|
|
promptTokens?: number;
|
|
completionTokens?: number;
|
|
response?: string;
|
|
}
|
|
|
|
function attachTrajectoryCapture(runtime: IAgentRuntime): CapturedLlmCall[] {
|
|
const calls: CapturedLlmCall[] = [];
|
|
const trajectoryLogger = {
|
|
isEnabled: () => true,
|
|
logLlmCall: (params: CapturedLlmCall) => {
|
|
calls.push(params);
|
|
},
|
|
};
|
|
const original = runtime.getServicesByType.bind(runtime);
|
|
runtime.getServicesByType = ((type: string) => {
|
|
if (type === "trajectories") return [trajectoryLogger];
|
|
return original(type);
|
|
}) as typeof runtime.getServicesByType;
|
|
const originalGet = runtime.getService.bind(runtime);
|
|
runtime.getService = ((name: string) => {
|
|
if (name === "trajectories") return trajectoryLogger;
|
|
return originalGet(name);
|
|
}) as typeof runtime.getService;
|
|
return calls;
|
|
}
|
|
|
|
describeLive(
|
|
"OpenAI trajectory wrapping (live)",
|
|
{ requiredEnv: ["OPENAI_API_KEY"] },
|
|
({ harness }) => {
|
|
it("records text and structured-output generation through recordLlmCall", async () => {
|
|
const { runtime } = harness();
|
|
const calls = attachTrajectoryCapture(runtime);
|
|
|
|
await runWithTrajectoryContext({ trajectoryStepId: "step-openai-live" }, async () => {
|
|
await handleTextSmall(runtime, {
|
|
prompt: "Reply with the single word: hello",
|
|
});
|
|
await handleTextLarge(runtime, {
|
|
prompt: 'Return JSON with shape {"ok": true} and nothing else.',
|
|
responseSchema: {
|
|
type: "object",
|
|
properties: { ok: { type: "boolean" } },
|
|
required: ["ok"],
|
|
},
|
|
} as Parameters<typeof handleTextLarge>[1]);
|
|
});
|
|
|
|
expect(calls).toHaveLength(2);
|
|
const [textCall, structuredCall] = calls;
|
|
expect(textCall.stepId).toBe("step-openai-live");
|
|
expect(textCall.actionType).toBe("ai.generateText");
|
|
expect(textCall.promptTokens ?? 0).toBeGreaterThan(0);
|
|
expect(textCall.completionTokens ?? 0).toBeGreaterThan(0);
|
|
expect(structuredCall.stepId).toBe("step-openai-live");
|
|
expect(structuredCall.actionType).toBe("ai.generateText");
|
|
expect(structuredCall.promptTokens ?? 0).toBeGreaterThan(0);
|
|
expect(structuredCall.completionTokens ?? 0).toBeGreaterThan(0);
|
|
}, 120_000);
|
|
}
|
|
);
|
|
|
|
// OpenAI-only paths (image generation, research, audio) are not supported by
|
|
// Cerebras. They are gated behind OPENAI_API_KEY_REAL so users can opt into
|
|
// real-OpenAI-only assertions without affecting the default Cerebras run.
|
|
const hasRealOpenAI =
|
|
Boolean(process.env.OPENAI_API_KEY_REAL?.trim()) &&
|
|
!process.env.OPENAI_BASE_URL?.includes("cerebras.ai");
|
|
|
|
describe.skipIf(!hasRealOpenAI)("OpenAI image/audio trajectory (real OpenAI)", () => {
|
|
it("records image, research, and audio generation calls", async () => {
|
|
const previousKey = process.env.OPENAI_API_KEY;
|
|
process.env.OPENAI_API_KEY = process.env.OPENAI_API_KEY_REAL;
|
|
try {
|
|
const { handleImageGeneration } = await import("../models/image");
|
|
const { handleTextToSpeech } = await import("../models/audio");
|
|
const calls: CapturedLlmCall[] = [];
|
|
const runtime = {
|
|
agentId: "agent-openai-real",
|
|
character: { system: "system prompt" },
|
|
emitEvent: () => {},
|
|
getService: (name: string) =>
|
|
name === "trajectories"
|
|
? {
|
|
isEnabled: () => true,
|
|
logLlmCall: (p: CapturedLlmCall) => calls.push(p),
|
|
}
|
|
: null,
|
|
getServicesByType: (type: string) =>
|
|
type === "trajectories"
|
|
? [
|
|
{
|
|
isEnabled: () => true,
|
|
logLlmCall: (p: CapturedLlmCall) => calls.push(p),
|
|
},
|
|
]
|
|
: [],
|
|
getSetting: (key: string) => process.env[key],
|
|
} as IAgentRuntime;
|
|
|
|
await runWithTrajectoryContext({ trajectoryStepId: "step-openai-real" }, async () => {
|
|
await handleImageGeneration(runtime, {
|
|
prompt: "A small red cube on a white background",
|
|
});
|
|
await handleTextToSpeech(runtime, "Hello live test");
|
|
});
|
|
|
|
expect(calls.map((c) => c.actionType)).toContain("openai.images.generate");
|
|
expect(calls.map((c) => c.actionType)).toContain("openai.audio.speech.create");
|
|
} finally {
|
|
if (previousKey === undefined) delete process.env.OPENAI_API_KEY;
|
|
else process.env.OPENAI_API_KEY = previousKey;
|
|
}
|
|
}, 180_000);
|
|
});
|