elizaos--eliza
426e9eeabd
Voice Workbench / headless workbench (mocked backends) (push) Has been cancelled
Voice Workbench / real acoustic lane (nightly, provisioned only) (push) Has been cancelled
ci / test (push) Has been cancelled
ci / lint-and-format (push) Has been cancelled
ci / build (push) Has been cancelled
ci / dev-startup (push) Has been cancelled
gitleaks / gitleaks (push) Has been cancelled
Markdown Links / Relative Markdown Links (push) Has been cancelled
Quality (Extended) / Homepage Build (PR smoke) (push) Has been cancelled
Quality (Extended) / Comment-only diff guard (push) Has been cancelled
Quality (Extended) / Format + Type Safety Ratchet (push) Has been cancelled
Quality (Extended) / Develop Gate (secret scan + UI determinism) (push) Has been cancelled
Quality (Extended) / Develop Gate (lint) (push) Has been cancelled
Chat shell gestures / Chat shell gesture + parity e2e (push) Has been cancelled
Cloud Gateway Discord / Test (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx @biomejs/biome check packages/lifeops-bench/src, benchmark-lint) (push) Has been cancelled
Benchmark Bridge Tests / benchmark (bunx vitest run --config packages/lifeops-bench/vitest.config.ts --root packages/lifeops-bench --passWithNoTests, benchmark-tests) (push) Has been cancelled
Build Agent Image / build-and-push (push) Has been cancelled
Dev Smoke / bun run dev onboarding chat (push) Has been cancelled
Dev Smoke / Vite HMR dependency-level smoke (push) Has been cancelled
Electrobun Submodule Guard / electrobun gitlink is fetchable (push) Has been cancelled
Publish @elizaos/example-code / check_npm (push) Has been cancelled
Publish @elizaos/example-code / publish_npm (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / verify_version (push) Has been cancelled
Publish @elizaos/plugin-elizacloud / publish_npm (push) Has been cancelled
Sandbox Live Smoke / Sandbox live smoke (push) Has been cancelled
Snap Build & Test / Build Snap (amd64) (push) Has been cancelled
Snap Build & Test / Build Snap (arm64) (push) Has been cancelled
Test Packaging / elizaos CLI global-install smoke (node + bun) (push) Has been cancelled
Cloud Gateway Webhook / Test (push) Has been cancelled
Cloud Tests / lint-and-types (push) Has been cancelled
Cloud Tests / unit-tests (push) Has been cancelled
Cloud Tests / integration-tests (push) Has been cancelled
Cloud Tests / e2e-tests (push) Has been cancelled
CodeQL Advanced / Analyze (javascript-typescript) (push) Has been cancelled
Deploy Apps Worker (Product 2) / Determine environment (push) Has been cancelled
Deploy Apps Worker (Product 2) / Deploy apps worker to apps-control host (${{ needs.determine-env.outputs.environment }}) (push) Has been cancelled
Deploy Eliza Provisioning Worker / Determine environment (push) Has been cancelled
Deploy Eliza Provisioning Worker / Deploy worker to Hetzner host (${{ needs.determine-env.outputs.environment }} @ ${{ needs.determine-env.outputs.deployment_sha }}) (push) Has been cancelled
Dev Smoke / Classify changed paths (push) Has been cancelled
supply-chain / sbom (push) Has been cancelled
supply-chain / vulnerability-scan (push) Has been cancelled
Build, Push & Deploy to Phala Cloud / build-and-push (push) Has been cancelled
Test Packaging / Validate Packaging Configs (push) Has been cancelled
Test Packaging / Build & Test PyPI Package (push) Has been cancelled
Test Packaging / PyPI on Python ${{ matrix.python }} (push) Has been cancelled
Test Packaging / Pack & Test JS Tarballs (push) Has been cancelled
UI Fixture E2E / ui-fixture-e2e (push) Has been cancelled
UI Fixture E2E / fixture-e2e (push) Has been cancelled
UI Story Gate / story-gate (push) Has been cancelled
vault-ci / test (macos-latest) (push) Has been cancelled
vault-ci / test (ubuntu-latest) (push) Has been cancelled
vault-ci / test (windows-latest) (push) Has been cancelled
vault-ci / app-core wiring tests (push) Has been cancelled
verify-patches / verify patches/CHECKSUMS.sha256 (push) Has been cancelled
Voice Benchmark Smoke / voice-emotion fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voiceagentbench fixture smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench-quality unit smoke (push) Has been cancelled
Voice Benchmark Smoke / voicebench TypeScript unit (no audio) (push) Has been cancelled
Voice Benchmark Smoke / voice bench smoke summary (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/app-core test bun run --cwd packages/elizaos test bun run --cwd packages/cloud/shared test], app-and-cli) (push) Has been cancelled
Windows CI / windows ([bun run --cwd packages/scenario-runner test bun run --cwd packages/vault test bun run --cwd packages/security test bun run --cwd plugins/plugin-coding-tools test], framework-packages) (push) Has been cancelled
Windows CI / windows ([bun run --cwd plugins/plugin-elizacloud test bun run --cwd plugins/plugin-discord test bun run --cwd plugins/plugin-anthropic test bun run --cwd plugins/plugin-openai test bun run --cwd plugins/plugin-app-control test bun run --cwd plugins/pl… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run build --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/agent --concurrency=4 node packages/scripts/run-bash-linux-only.mjs scripts/verify-riscv64-buildpaths.sh node packages/scripts/run… (push) Has been cancelled
Windows CI / windows ([node packages/scripts/run-turbo.mjs run typecheck --filter=@elizaos/core --filter=@elizaos/shared --filter=@elizaos/cloud-shared --concurrency=4 bun run --cwd packages/core test bun run --cwd packages/shared test], core-runtime, 75) (push) Has been cancelled
157 行
5.4 KiB
TypeScript
157 行
5.4 KiB
TypeScript
/**
|
|
* Live test that real `TEXT_SMALL`/`TEXT_LARGE` and native tool-call generation
|
|
* flow through `recordLlmCall` into the trajectory logger with correct step id,
|
|
* action type, token counts, and response. Hits the real Gemini API and
|
|
* self-skips when `GOOGLE_GENERATIVE_AI_API_KEY` is unset.
|
|
*/
|
|
import type { IAgentRuntime } from "@elizaos/core";
|
|
import { runWithTrajectoryContext } from "@elizaos/core";
|
|
import { describe, expect, it } from "vitest";
|
|
|
|
interface CapturedLlmCall {
|
|
stepId: string;
|
|
actionType: string;
|
|
promptTokens?: number;
|
|
completionTokens?: number;
|
|
response?: string;
|
|
}
|
|
|
|
const REQUIRED_KEY = "GOOGLE_GENERATIVE_AI_API_KEY";
|
|
const apiKey = process.env[REQUIRED_KEY]?.trim();
|
|
const SHOULD_RUN = Boolean(apiKey);
|
|
|
|
function createInlineRuntime(calls: CapturedLlmCall[]): IAgentRuntime {
|
|
const trajectoryLogger = {
|
|
isEnabled: () => true,
|
|
logLlmCall: (params: CapturedLlmCall) => {
|
|
calls.push(params);
|
|
},
|
|
};
|
|
const settings: Record<string, string> = {
|
|
GOOGLE_GENERATIVE_AI_API_KEY: apiKey ?? "",
|
|
};
|
|
return {
|
|
agentId: "agent-google",
|
|
character: { system: "You are a concise assistant." },
|
|
emitEvent: async () => undefined,
|
|
getService: (name: string) =>
|
|
name === "trajectories" ? trajectoryLogger : null,
|
|
getServicesByType: (type: string) =>
|
|
type === "trajectories" ? [trajectoryLogger] : [],
|
|
getSetting: (key: string) => settings[key] ?? process.env[key] ?? null,
|
|
} as IAgentRuntime;
|
|
}
|
|
|
|
if (!SHOULD_RUN) {
|
|
process.env.SKIP_REASON ||= `missing required env: ${REQUIRED_KEY}`;
|
|
console.warn(
|
|
`\x1b[33m[google-genai trajectory.test] live test disabled: missing required env ${REQUIRED_KEY} (set ${REQUIRED_KEY} to enable)\x1b[0m`,
|
|
);
|
|
describe("Google GenAI trajectory wrapping (live)", () => {
|
|
it.skip(`[live] requires ${REQUIRED_KEY}`, () => {});
|
|
});
|
|
} else {
|
|
describe("Google GenAI trajectory wrapping (live)", () => {
|
|
it("records text and structured-output generation via TEXT_* through recordLlmCall", async () => {
|
|
const { handleTextSmall, handleTextLarge } = await import(
|
|
"../models/text"
|
|
);
|
|
|
|
const calls: CapturedLlmCall[] = [];
|
|
const runtime = createInlineRuntime(calls);
|
|
|
|
await runWithTrajectoryContext(
|
|
{ trajectoryStepId: "step-google" },
|
|
async () => {
|
|
await handleTextSmall(runtime, {
|
|
prompt: "What is 2+2? Reply with just the number.",
|
|
maxTokens: 32,
|
|
});
|
|
await handleTextLarge(runtime, {
|
|
prompt:
|
|
'Return JSON {"answer": 4} for the question 2+2. Reply with only the JSON object.',
|
|
responseSchema: {
|
|
type: "object",
|
|
properties: { answer: { type: "number" } },
|
|
required: ["answer"],
|
|
},
|
|
} as Parameters<typeof handleTextLarge>[1]);
|
|
},
|
|
);
|
|
|
|
expect(calls).toHaveLength(2);
|
|
const [textCall, structuredCall] = calls;
|
|
expect(textCall.stepId).toBe("step-google");
|
|
expect(textCall.actionType).toBe(
|
|
"google-genai.TEXT_SMALL.generateContent",
|
|
);
|
|
expect(textCall.promptTokens ?? 0).toBeGreaterThan(0);
|
|
expect(textCall.completionTokens ?? 0).toBeGreaterThan(0);
|
|
expect(textCall.response).toContain("4");
|
|
expect(structuredCall.stepId).toBe("step-google");
|
|
expect(structuredCall.actionType).toBe(
|
|
"google-genai.TEXT_LARGE.generateContent",
|
|
);
|
|
expect(structuredCall.promptTokens ?? 0).toBeGreaterThan(0);
|
|
expect(structuredCall.completionTokens ?? 0).toBeGreaterThan(0);
|
|
expect(structuredCall.response).toContain("4");
|
|
}, 120_000);
|
|
|
|
it("records a native Gemini tool call trajectory", async () => {
|
|
const { handleTextSmall } = await import("../models/text");
|
|
|
|
const calls: CapturedLlmCall[] = [];
|
|
const runtime = createInlineRuntime(calls);
|
|
|
|
const result = (await runWithTrajectoryContext(
|
|
{ trajectoryStepId: "step-google-tool-call" },
|
|
async () =>
|
|
handleTextSmall(runtime, {
|
|
prompt:
|
|
"Use the lookup_weather tool for Paris. Do not answer in plain text.",
|
|
maxTokens: 128,
|
|
tools: {
|
|
lookup_weather: {
|
|
description: "Lookup current weather for a city.",
|
|
inputSchema: {
|
|
type: "object",
|
|
properties: {
|
|
city: { type: "string" },
|
|
},
|
|
required: ["city"],
|
|
},
|
|
},
|
|
},
|
|
toolChoice: { type: "tool", toolName: "lookup_weather" },
|
|
} as Parameters<typeof handleTextSmall>[1]),
|
|
)) as unknown as {
|
|
text: string;
|
|
toolCalls?: Array<{
|
|
name?: string;
|
|
toolName?: string;
|
|
input?: unknown;
|
|
arguments?: unknown;
|
|
}>;
|
|
finishReason?: string;
|
|
};
|
|
|
|
expect(result.toolCalls?.length ?? 0).toBeGreaterThan(0);
|
|
expect(
|
|
result.toolCalls?.some(
|
|
(call) =>
|
|
call.name === "lookup_weather" ||
|
|
call.toolName === "lookup_weather",
|
|
),
|
|
).toBe(true);
|
|
expect(result.finishReason).toBe("tool-calls");
|
|
|
|
expect(calls).toHaveLength(1);
|
|
expect(calls[0]?.stepId).toBe("step-google-tool-call");
|
|
expect(calls[0]?.actionType).toBe(
|
|
"google-genai.TEXT_SMALL.generateContent",
|
|
);
|
|
expect(calls[0]?.response ?? result.text).toBe(result.text);
|
|
}, 120_000);
|
|
});
|
|
}
|