Files
wehub-resource-sync 426e9eeabd
Voice Workbench / headless workbench (mocked backends) (push) Has been cancelled
Voice Workbench / real acoustic lane (nightly, provisioned only) (push) Has been cancelled
ci / test (push) Has been cancelled
ci / lint-and-format (push) Has been cancelled
ci / build (push) Has been cancelled
ci / dev-startup (push) Has been cancelled
gitleaks / gitleaks (push) Has been cancelled
Markdown Links / Relative Markdown Links (push) Has been cancelled
Quality (Extended) / Homepage Build (PR smoke) (push) Has been cancelled
Quality (Extended) / Comment-only diff guard (push) Has been cancelled
Quality (Extended) / Format + Type Safety Ratchet (push) Has been cancelled
Quality (Extended) / Develop Gate (secret scan + UI determinism) (push) Has been cancelled
Quality (Extended) / Develop Gate (lint) (push) Has been cancelled
Chat shell gestures / Chat shell gesture + parity e2e (push) Has been cancelled
Cloud Gateway Discord / Test (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx @biomejs/biome check packages/lifeops-bench/src, benchmark-lint) (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx vitest run --config packages/lifeops-bench/vitest.config.ts --root packages/lifeops-bench --passWithNoTests, benchmark-tests) (push) Has been cancelled
Build Agent Image / build-and-push (push) Has been cancelled
Dev Smoke / bun run dev onboarding chat (push) Has been cancelled
Dev Smoke / Vite HMR dependency-level smoke (push) Has been cancelled
Electrobun Submodule Guard / electrobun gitlink is fetchable (push) Has been cancelled
Publish @elizaos/example-code / check_npm (push) Has been cancelled
Publish @elizaos/example-code / publish_npm (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / verify_version (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / publish_npm (push) Has been cancelled
Sandbox Live Smoke / Sandbox live smoke (push) Has been cancelled
Snap Build & Test / Build Snap (amd64) (push) Has been cancelled
Snap Build & Test / Build Snap (arm64) (push) Has been cancelled
Test Packaging / elizaos CLI global-install smoke (node + bun) (push) Has been cancelled
Cloud Gateway Webhook / Test (push) Has been cancelled
Cloud Tests / lint-and-types (push) Has been cancelled
Cloud Tests / unit-tests (push) Has been cancelled
Cloud Tests / integration-tests (push) Has been cancelled
Cloud Tests / e2e-tests (push) Has been cancelled
CodeQL Advanced / Analyze (javascript-typescript) (push) Has been cancelled
Deploy Apps Worker (Product 2) / Determine environment (push) Has been cancelled
Deploy Apps Worker (Product 2) / Deploy apps worker to apps-control host (${{ needs.determine-env.outputs.environment }}) (push) Has been cancelled
Deploy Eliza Provisioning Worker / Determine environment (push) Has been cancelled
Deploy Eliza Provisioning Worker / Deploy worker to Hetzner host (${{ needs.determine-env.outputs.environment }} @ ${{ needs.determine-env.outputs.deployment_sha }}) (push) Has been cancelled
Dev Smoke / Classify changed paths (push) Has been cancelled
supply-chain / sbom (push) Has been cancelled
supply-chain / vulnerability-scan (push) Has been cancelled
Build, Push & Deploy to Phala Cloud / build-and-push (push) Has been cancelled
Test Packaging / Validate Packaging Configs (push) Has been cancelled
Test Packaging / Build & Test PyPI Package (push) Has been cancelled
Test Packaging / PyPI on Python ${{ matrix.python }} (push) Has been cancelled
Test Packaging / Pack & Test JS Tarballs (push) Has been cancelled
UI Fixture E2E / ui-fixture-e2e (push) Has been cancelled
UI Fixture E2E / fixture-e2e (push) Has been cancelled
UI Story Gate / story-gate (push) Has been cancelled
vault-ci / test (macos-latest) (push) Has been cancelled
vault-ci / test (ubuntu-latest) (push) Has been cancelled
vault-ci / test (windows-latest) (push) Has been cancelled
vault-ci / app-core wiring tests (push) Has been cancelled
verify-patches / verify patches/CHECKSUMS.sha256 (push) Has been cancelled
Voice Benchmark Smoke / voice-emotion fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voiceagentbench fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench-quality unit smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench TypeScript unit (no audio) (push) Has been cancelled
Voice Benchmark Smoke / voice bench smoke summary (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/app-core test bun run --cwd packages/elizaos test bun run --cwd packages/cloud/shared test], app-and-cli) (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/scenario-runner test bun run --cwd packages/vault test bun run --cwd packages/security test bun run --cwd plugins/plugin-coding-tools test], framework-packages) (push) Has been cancelled
Windows CI / windows ([bun run --cwd plugins/plugin-elizacloud test bun run --cwd plugins/plugin-discord test bun run --cwd plugins/plugin-anthropic test bun run --cwd plugins/plugin-openai test bun run --cwd plugins/plugin-app-control test bun run --cwd plugins/pl… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run build --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/agent --concurrency=4 node packages/scripts/run-bash-linux-only.mjs scripts/verify-riscv64-buildpaths.sh node packages/scripts/run… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run typecheck --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/cloud-shared --concurrency=4 bun run --cwd packages/core test bun run --cwd packages/shared test], core-runtime, 75) (push) Has been cancelled
chore: import upstream snapshot with attribution
2026-07-13 12:43:05 +08:00

130 lines
4.9 KiB
TypeScript

/**
* Live test asserting that the text handlers record their LLM call into the
* active trajectory context. Post-merge lane, real model.
*/
import type { IAgentRuntime } from "@elizaos/core";
import { runWithTrajectoryContext } from "@elizaos/core";
import { describe, expect, it } from "vitest";
import { describeLive } from "../../../packages/app-core/test/helpers/live-agent-test";
import { handleTextLarge, handleTextSmall } from "../models/text";
interface CapturedLlmCall {
stepId: string;
actionType: string;
promptTokens?: number;
completionTokens?: number;
response?: string;
}
function attachTrajectoryCapture(runtime: IAgentRuntime): CapturedLlmCall[] {
const calls: CapturedLlmCall[] = [];
const trajectoryLogger = {
isEnabled: () => true,
logLlmCall: (params: CapturedLlmCall) => {
calls.push(params);
},
};
const original = runtime.getServicesByType.bind(runtime);
runtime.getServicesByType = ((type: string) => {
if (type === "trajectories") return [trajectoryLogger];
return original(type);
}) as typeof runtime.getServicesByType;
const originalGet = runtime.getService.bind(runtime);
runtime.getService = ((name: string) => {
if (name === "trajectories") return trajectoryLogger;
return originalGet(name);
}) as typeof runtime.getService;
return calls;
}
describeLive(
"OpenAI trajectory wrapping (live)",
{ requiredEnv: ["OPENAI_API_KEY"] },
({ harness }) => {
it("records text and structured-output generation through recordLlmCall", async () => {
const { runtime } = harness();
const calls = attachTrajectoryCapture(runtime);
await runWithTrajectoryContext({ trajectoryStepId: "step-openai-live" }, async () => {
await handleTextSmall(runtime, {
prompt: "Reply with the single word: hello",
});
await handleTextLarge(runtime, {
prompt: 'Return JSON with shape {"ok": true} and nothing else.',
responseSchema: {
type: "object",
properties: { ok: { type: "boolean" } },
required: ["ok"],
},
} as Parameters<typeof handleTextLarge>[1]);
});
expect(calls).toHaveLength(2);
const [textCall, structuredCall] = calls;
expect(textCall.stepId).toBe("step-openai-live");
expect(textCall.actionType).toBe("ai.generateText");
expect(textCall.promptTokens ?? 0).toBeGreaterThan(0);
expect(textCall.completionTokens ?? 0).toBeGreaterThan(0);
expect(structuredCall.stepId).toBe("step-openai-live");
expect(structuredCall.actionType).toBe("ai.generateText");
expect(structuredCall.promptTokens ?? 0).toBeGreaterThan(0);
expect(structuredCall.completionTokens ?? 0).toBeGreaterThan(0);
}, 120_000);
}
);
// OpenAI-only paths (image generation, research, audio) are not supported by
// Cerebras. They are gated behind OPENAI_API_KEY_REAL so users can opt into
// real-OpenAI-only assertions without affecting the default Cerebras run.
const hasRealOpenAI =
Boolean(process.env.OPENAI_API_KEY_REAL?.trim()) &&
!process.env.OPENAI_BASE_URL?.includes("cerebras.ai");
describe.skipIf(!hasRealOpenAI)("OpenAI image/audio trajectory (real OpenAI)", () => {
it("records image, research, and audio generation calls", async () => {
const previousKey = process.env.OPENAI_API_KEY;
process.env.OPENAI_API_KEY = process.env.OPENAI_API_KEY_REAL;
try {
const { handleImageGeneration } = await import("../models/image");
const { handleTextToSpeech } = await import("../models/audio");
const calls: CapturedLlmCall[] = [];
const runtime = {
agentId: "agent-openai-real",
character: { system: "system prompt" },
emitEvent: () => {},
getService: (name: string) =>
name === "trajectories"
? {
isEnabled: () => true,
logLlmCall: (p: CapturedLlmCall) => calls.push(p),
}
: null,
getServicesByType: (type: string) =>
type === "trajectories"
? [
{
isEnabled: () => true,
logLlmCall: (p: CapturedLlmCall) => calls.push(p),
},
]
: [],
getSetting: (key: string) => process.env[key],
} as IAgentRuntime;
await runWithTrajectoryContext({ trajectoryStepId: "step-openai-real" }, async () => {
await handleImageGeneration(runtime, {
prompt: "A small red cube on a white background",
});
await handleTextToSpeech(runtime, "Hello live test");
});
expect(calls.map((c) => c.actionType)).toContain("openai.images.generate");
expect(calls.map((c) => c.actionType)).toContain("openai.audio.speech.create");
} finally {
if (previousKey === undefined) delete process.env.OPENAI_API_KEY;
else process.env.OPENAI_API_KEY = previousKey;
}
}, 180_000);
});