Some checks failed
ClawSweeper Dispatch / dispatch (push) Has been cancelled
CodeQL / Security High (actions) (push) Has been cancelled
CodeQL / Security High (channel-runtime-boundary) (push) Has been cancelled
CodeQL / Security High (core-auth-secrets) (push) Has been cancelled
CodeQL / Security High (mcp-process-tool-boundary) (push) Has been cancelled
CodeQL / Security High (network-ssrf-boundary) (push) Has been cancelled
CodeQL / Security High (plugin-trust-boundary) (push) Has been cancelled
CodeQL / Security High (process-exec-boundary) (push) Has been cancelled
Docs Sync Publish Repo / sync-publish-repo (push) Has been cancelled
Docs / docs (push) Has been cancelled
OpenClaw Stable Main Closeout / Resolve stable release closeout inputs (push) Has been cancelled
OpenClaw Stable Main Closeout / Verify stable main closeout (push) Has been cancelled
Workflow Sanity / no-tabs (push) Has been cancelled
Workflow Sanity / actionlint (push) Has been cancelled
Workflow Sanity / generated-doc-baselines (push) Has been cancelled
CI / runner-admission (push) Has been cancelled
CI / preflight (push) Has been cancelled
CI / security-fast (push) Has been cancelled
CI / pnpm-store-warmup (push) Has been cancelled
CI / build-artifacts (push) Has been cancelled
CI / native-i18n (push) Has been cancelled
CI / ${{ matrix.check_name }} (push) Has been cancelled
CI / ${{ matrix.checkName }} (push) Has been cancelled
CI / checks-node-compat-node22 (push) Has been cancelled
CI / check-bundled-channel-config-metadata (push) Has been cancelled
CI / check-dependencies (push) Has been cancelled
CI / check-guards (push) Has been cancelled
CI / check-lint (push) Has been cancelled
CI / check-prod-types (push) Has been cancelled
CI / check-shrinkwrap (push) Has been cancelled
CI / check-test-types (push) Has been cancelled
CI / check-additional-boundaries-a (push) Has been cancelled
CI / check-additional-boundaries-bcd (push) Has been cancelled
CI / check-additional-extension-bundled (push) Has been cancelled
CI / check-additional-extension-channels (push) Has been cancelled
CI / check-additional-extension-package-boundary (push) Has been cancelled
CI / check-additional-runtime-topology-architecture (push) Has been cancelled
CI / check-session-accessor-boundary (push) Has been cancelled
CI / check-session-transcript-reader-boundary (push) Has been cancelled
CI / check-docs (push) Has been cancelled
CI / skills-python (push) Has been cancelled
CI / macos-swift (push) Has been cancelled
CI / ios-build (push) Has been cancelled
CI / ci-timings-summary (push) Has been cancelled
Native App Locale Refresh / Refresh native fa (push) Has been cancelled
Native App Locale Refresh / Refresh native fr (push) Has been cancelled
Native App Locale Refresh / Refresh native hi (push) Has been cancelled
Native App Locale Refresh / Refresh native id (push) Has been cancelled
Native App Locale Refresh / Refresh native it (push) Has been cancelled
Native App Locale Refresh / Refresh native ja-JP (push) Has been cancelled
Control UI Locale Refresh / plan (push) Has been cancelled
Control UI Locale Refresh / Refresh ${{ matrix.locale }} (push) Has been cancelled
Control UI Locale Refresh / Commit control UI locale refresh (push) Has been cancelled
Live Media Runner Image / Build live media runner image (push) Has been cancelled
Native App Locale Refresh / Refresh native ar (push) Has been cancelled
Native App Locale Refresh / Refresh native de (push) Has been cancelled
Native App Locale Refresh / Refresh native es (push) Has been cancelled
Native App Locale Refresh / Refresh native ko (push) Has been cancelled
Native App Locale Refresh / Refresh native nl (push) Has been cancelled
Native App Locale Refresh / Refresh native pl (push) Has been cancelled
Native App Locale Refresh / Refresh native pt-BR (push) Has been cancelled
Native App Locale Refresh / Refresh native ru (push) Has been cancelled
Native App Locale Refresh / Refresh native sv (push) Has been cancelled
Native App Locale Refresh / Refresh native th (push) Has been cancelled
Native App Locale Refresh / Refresh native tr (push) Has been cancelled
Native App Locale Refresh / Refresh native uk (push) Has been cancelled
Native App Locale Refresh / Refresh native vi (push) Has been cancelled
Native App Locale Refresh / Refresh native zh-CN (push) Has been cancelled
Native App Locale Refresh / Refresh native zh-TW (push) Has been cancelled
Native App Locale Refresh / Commit native locale refresh (push) Has been cancelled
Plugin Init Scaffold Validation / Validate provider scaffold (push) Has been cancelled
Plugin NPM Release / preview_plugins_npm (push) Has been cancelled
Plugin NPM Release / Validate release publish approval (push) Has been cancelled
Plugin NPM Release / preview_plugin_pack (push) Has been cancelled
Plugin NPM Release / publish_plugins_npm (push) Has been cancelled
Sandbox Common Smoke / sandbox-common-smoke (push) Has been cancelled
Website Installer Sync / static (push) Has been cancelled
Website Installer Sync / linux-docker (push) Has been cancelled
Website Installer Sync / macos-installer (push) Has been cancelled
Website Installer Sync / windows-installer (push) Has been cancelled
Website Installer Sync / sync-website (push) Has been cancelled
Adolf is a fork/vendored clone of github.com/openclaw/openclaw (v2026.6.11), free to diverge. Tree copied sans upstream .git; upstream remote added for future syncs. Node pinned to 24 (.nvmrc); engines already require >=22.19. Preserves docs/ARCHITECTURE.md. Co-Authored-By: Claude Opus 4.8 <noreply@anthropic.com> Claude-Session: https://claude.ai/code/session_01LeqyaxJF2nbRXJtae2kNB2
331 lines
9.3 KiB
TypeScript
331 lines
9.3 KiB
TypeScript
// Vllm tests cover stream plugin behavior.
|
|
import type { StreamFn } from "openclaw/plugin-sdk/agent-core";
|
|
import type { Context, Model } from "openclaw/plugin-sdk/llm";
|
|
import { describe, expect, it } from "vitest";
|
|
import {
|
|
createVllmProviderThinkingWrapper,
|
|
createVllmQwenThinkingWrapper,
|
|
wrapVllmProviderStream,
|
|
} from "./stream.js";
|
|
|
|
function capturePayload(params: {
|
|
format: "chat-template" | "top-level";
|
|
thinkingLevel?: "off" | "low" | "medium" | "high" | "xhigh" | "max";
|
|
reasoning?: unknown;
|
|
initialPayload?: Record<string, unknown>;
|
|
model?: Partial<Model<"openai-completions">>;
|
|
}): Record<string, unknown> {
|
|
let captured: Record<string, unknown> = {};
|
|
const baseStreamFn: StreamFn = (_model, _context, options) => {
|
|
const payload = { ...params.initialPayload };
|
|
options?.onPayload?.(payload, _model);
|
|
captured = payload;
|
|
return {} as ReturnType<StreamFn>;
|
|
};
|
|
|
|
const wrapped = createVllmQwenThinkingWrapper({
|
|
baseStreamFn,
|
|
format: params.format,
|
|
thinkingLevel: params.thinkingLevel ?? "high",
|
|
});
|
|
void wrapped(
|
|
{
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "Qwen/Qwen3-8B",
|
|
reasoning: true,
|
|
...params.model,
|
|
} as Model<"openai-completions">,
|
|
{ messages: [] } as Context,
|
|
params.reasoning === undefined ? {} : ({ reasoning: params.reasoning } as never),
|
|
);
|
|
|
|
return captured;
|
|
}
|
|
|
|
describe("createVllmQwenThinkingWrapper", () => {
|
|
it("maps Qwen chat-template thinking off to chat_template_kwargs", () => {
|
|
const payload = capturePayload({
|
|
format: "chat-template",
|
|
reasoning: "none",
|
|
initialPayload: {
|
|
reasoning_effort: "high",
|
|
reasoning: { effort: "high" },
|
|
reasoningEffort: "high",
|
|
},
|
|
});
|
|
|
|
expect(payload).toEqual({
|
|
chat_template_kwargs: {
|
|
enable_thinking: false,
|
|
preserve_thinking: true,
|
|
},
|
|
});
|
|
});
|
|
|
|
it("maps Qwen chat-template thinking on to chat_template_kwargs", () => {
|
|
expect(capturePayload({ format: "chat-template", reasoning: "medium" })).toEqual({
|
|
chat_template_kwargs: {
|
|
enable_thinking: true,
|
|
preserve_thinking: true,
|
|
},
|
|
});
|
|
});
|
|
|
|
it("preserves explicit chat-template kwargs while setting enable_thinking", () => {
|
|
expect(
|
|
capturePayload({
|
|
format: "chat-template",
|
|
thinkingLevel: "off",
|
|
initialPayload: {
|
|
chat_template_kwargs: {
|
|
preserve_thinking: false,
|
|
force_nonempty_content: true,
|
|
},
|
|
},
|
|
}),
|
|
).toEqual({
|
|
chat_template_kwargs: {
|
|
enable_thinking: false,
|
|
preserve_thinking: false,
|
|
force_nonempty_content: true,
|
|
},
|
|
});
|
|
});
|
|
|
|
it("maps Qwen top-level thinking format to enable_thinking", () => {
|
|
expect(capturePayload({ format: "top-level", thinkingLevel: "off" })).toEqual({
|
|
enable_thinking: false,
|
|
});
|
|
expect(capturePayload({ format: "top-level", thinkingLevel: "high" })).toEqual({
|
|
enable_thinking: true,
|
|
});
|
|
});
|
|
|
|
it("patches configured Qwen models unless reasoning is explicitly disabled", () => {
|
|
expect(capturePayload({ format: "chat-template", model: { reasoning: undefined } })).toEqual({
|
|
chat_template_kwargs: {
|
|
enable_thinking: true,
|
|
preserve_thinking: true,
|
|
},
|
|
});
|
|
expect(capturePayload({ format: "chat-template", model: { reasoning: false } })).toStrictEqual(
|
|
{},
|
|
);
|
|
});
|
|
|
|
it("skips non-completions models", () => {
|
|
expect(
|
|
capturePayload({ format: "chat-template", model: { api: "openai-responses" as never } }),
|
|
).toStrictEqual({});
|
|
});
|
|
});
|
|
|
|
describe("createVllmProviderThinkingWrapper", () => {
|
|
function captureProviderPayload(params: {
|
|
thinkingLevel?: "off" | "low" | "medium" | "high" | "xhigh" | "max";
|
|
initialPayload?: Record<string, unknown>;
|
|
model?: Partial<Model<"openai-completions">>;
|
|
}): Record<string, unknown> {
|
|
let captured: Record<string, unknown> = {};
|
|
const baseStreamFn: StreamFn = (_model, _context, options) => {
|
|
const payload = { ...params.initialPayload };
|
|
options?.onPayload?.(payload, _model);
|
|
captured = payload;
|
|
return {} as ReturnType<StreamFn>;
|
|
};
|
|
|
|
const wrapped = createVllmProviderThinkingWrapper({
|
|
baseStreamFn,
|
|
thinkingLevel: params.thinkingLevel ?? "high",
|
|
});
|
|
void wrapped(
|
|
{
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "nemotron-3-super",
|
|
reasoning: true,
|
|
...params.model,
|
|
} as Model<"openai-completions">,
|
|
{ messages: [] } as Context,
|
|
{},
|
|
);
|
|
|
|
return captured;
|
|
}
|
|
|
|
it("injects Nemotron 3 chat-template kwargs when thinking is off", () => {
|
|
expect(captureProviderPayload({ thinkingLevel: "off" })).toEqual({
|
|
chat_template_kwargs: {
|
|
enable_thinking: false,
|
|
force_nonempty_content: true,
|
|
},
|
|
});
|
|
});
|
|
|
|
it("does not inject Nemotron 3 chat-template kwargs when thinking is enabled", () => {
|
|
expect(captureProviderPayload({ thinkingLevel: "low" })).toStrictEqual({});
|
|
});
|
|
|
|
it("preserves existing Nemotron 3 chat-template kwargs over defaults", () => {
|
|
expect(
|
|
captureProviderPayload({
|
|
thinkingLevel: "off",
|
|
initialPayload: {
|
|
chat_template_kwargs: {
|
|
enable_thinking: true,
|
|
},
|
|
},
|
|
}),
|
|
).toEqual({
|
|
chat_template_kwargs: {
|
|
enable_thinking: true,
|
|
force_nonempty_content: true,
|
|
},
|
|
});
|
|
});
|
|
|
|
it("skips non-Nemotron vLLM models", () => {
|
|
expect(
|
|
captureProviderPayload({
|
|
thinkingLevel: "off",
|
|
model: { id: "Qwen/Qwen3-8B" },
|
|
}),
|
|
).toStrictEqual({});
|
|
});
|
|
});
|
|
|
|
describe("wrapVllmProviderStream", () => {
|
|
it("registers when vLLM Qwen thinking format compat is configured", () => {
|
|
expect(
|
|
wrapVllmProviderStream({
|
|
provider: "vllm",
|
|
modelId: "Qwen/Qwen3-8B",
|
|
extraParams: {},
|
|
model: {
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "Qwen/Qwen3-8B",
|
|
reasoning: true,
|
|
compat: { thinkingFormat: "qwen-chat-template" },
|
|
} as Model<"openai-completions">,
|
|
streamFn: undefined,
|
|
} as never),
|
|
).toBeTypeOf("function");
|
|
});
|
|
|
|
it("ignores request params when Qwen thinking format compat is not configured", () => {
|
|
expect(
|
|
wrapVllmProviderStream({
|
|
provider: "vllm",
|
|
modelId: "Qwen/Qwen3-8B",
|
|
extraParams: { qwenThinkingFormat: "chat-template" },
|
|
model: {
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "Qwen/Qwen3-8B",
|
|
reasoning: true,
|
|
} as Model<"openai-completions">,
|
|
streamFn: undefined,
|
|
} as never),
|
|
).toBeUndefined();
|
|
});
|
|
|
|
it("uses model compat for Qwen thinking format", () => {
|
|
let captured: Record<string, unknown> = {};
|
|
const baseStreamFn: StreamFn = (_model, _context, options) => {
|
|
const payload = {};
|
|
options?.onPayload?.(payload, _model);
|
|
captured = payload;
|
|
return {} as ReturnType<StreamFn>;
|
|
};
|
|
const model = {
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "Qwen/Qwen3-8B",
|
|
reasoning: true,
|
|
compat: { thinkingFormat: "qwen-chat-template" },
|
|
} as unknown as Model<"openai-completions">;
|
|
const wrapped = wrapVllmProviderStream({
|
|
provider: "vllm",
|
|
modelId: "Qwen/Qwen3-8B",
|
|
extraParams: {},
|
|
thinkingLevel: "off",
|
|
model,
|
|
streamFn: baseStreamFn,
|
|
} as never);
|
|
|
|
expect(wrapped).toBeTypeOf("function");
|
|
void wrapped?.(model, { messages: [] } as Context, {});
|
|
|
|
expect(captured).toEqual({
|
|
chat_template_kwargs: {
|
|
enable_thinking: false,
|
|
preserve_thinking: true,
|
|
},
|
|
});
|
|
});
|
|
|
|
it("skips unconfigured vLLM and non-vLLM providers", () => {
|
|
expect(
|
|
wrapVllmProviderStream({
|
|
provider: "vllm",
|
|
modelId: "Qwen/Qwen3-8B",
|
|
extraParams: {},
|
|
model: {
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "Qwen/Qwen3-8B",
|
|
} as Model<"openai-completions">,
|
|
streamFn: undefined,
|
|
} as never),
|
|
).toBeUndefined();
|
|
|
|
expect(
|
|
wrapVllmProviderStream({
|
|
provider: "openai",
|
|
modelId: "gpt-5.4",
|
|
extraParams: {},
|
|
model: {
|
|
api: "openai-completions",
|
|
provider: "openai",
|
|
id: "gpt-5.4",
|
|
} as Model<"openai-completions">,
|
|
streamFn: undefined,
|
|
} as never),
|
|
).toBeUndefined();
|
|
});
|
|
|
|
it("registers for vLLM Nemotron when thinking is off", () => {
|
|
expect(
|
|
wrapVllmProviderStream({
|
|
provider: "vllm",
|
|
modelId: "nemotron-3-super",
|
|
extraParams: {},
|
|
thinkingLevel: "off",
|
|
model: {
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "nemotron-3-super",
|
|
} as Model<"openai-completions">,
|
|
streamFn: undefined,
|
|
} as never),
|
|
).toBeTypeOf("function");
|
|
|
|
expect(
|
|
wrapVllmProviderStream({
|
|
provider: "vllm",
|
|
modelId: "nemotron-3-super",
|
|
extraParams: {},
|
|
thinkingLevel: "low",
|
|
model: {
|
|
api: "openai-completions",
|
|
provider: "vllm",
|
|
id: "nemotron-3-super",
|
|
} as Model<"openai-completions">,
|
|
streamFn: undefined,
|
|
} as never),
|
|
).toBeUndefined();
|
|
});
|
|
});
|