* webui: Move static build output from `tools/server/public` to `build/ui` directory * refactor: Move to `tools/ui` * refactor: rename CMake variables and preprocessor defines - Rename LLAMA_BUILD_WEBUI -> LLAMA_BUILD_UI (old kept as deprecated) - Rename LLAMA_USE_PREBUILT_WEBUI -> LLAMA_USE_PREBUILT_UI (old kept as deprecated) - Backward compat: old vars auto-forward to new ones with DEPRECATION warning - Rename internal vars: WEBUI_SOURCE -> UI_SOURCE, WEBUI_SOURCE_DIR -> UI_SOURCE_DIR, etc. - Rename HF bucket: LLAMA_WEBUI_HF_BUCKET -> LLAMA_UI_HF_BUCKET - Emit both LLAMA_BUILD_WEBUI and LLAMA_BUILD_UI preprocessor defines - Emit both LLAMA_WEBUI_DEFAULT_ENABLED and LLAMA_UI_DEFAULT_ENABLED * refactor: rename CLI flags (--webui -> --ui) with backward compat - Add --ui/--no-ui (old --webui/--no-webui kept as deprecated aliases) - Add --ui-config (old --webui-config kept as deprecated alias) - Add --ui-config-file (old --webui-config-file kept as deprecated alias) - Add --ui-mcp-proxy/--no-ui-mcp-proxy (old --webui-mcp-proxy kept as deprecated) - Add new env vars: LLAMA_ARG_UI, LLAMA_ARG_UI_CONFIG, LLAMA_ARG_UI_CONFIG_FILE, LLAMA_ARG_UI_MCP_PROXY - C++ struct fields: params.ui, params.ui_config_json, params.ui_mcp_proxy added alongside old fields - Backward compat: old fields synced to new ones in g_params_to_internals * refactor: update C++ server internals with backward compat - Rename json_webui_settings -> json_ui_settings (both kept in server_context_meta) - Rename params.webui usage -> params.ui (both synced, old still works) - JSON API emits both "ui"/"ui_settings" and "webui"/"webui_settings" keys - Server routes use params.ui_mcp_proxy || params.webui_mcp_proxy - Preprocessor guards use #if defined(LLAMA_BUILD_UI) || defined(LLAMA_BUILD_WEBUI) * refactor: rename CI/CD workflows, artifacts, and build script - Rename webui-build.yml -> ui-build.yml; artifact webui-build -> ui-build - Rename webui-publish.yml -> ui-publish.yml; var HF_BUCKET_WEBUI_STATIC_OUTPUT -> HF_BUCKET_UI_STATIC_OUTPUT - Rename server-webui.yml -> server-ui.yml; job webui-build/checks -> ui-build/checks - Update server.yml: job/artifact refs webui-build -> ui-build - Update release.yml: all webui-build/publish refs -> ui-build/publish; HF_TOKEN_WEBUI_STATIC_OUTPUT -> HF_TOKEN_UI_STATIC_OUTPUT - Update server-self-hosted.yml: webui-build -> ui-build - Update build-self-hosted.yml: HF_WEBUI_VERSION -> HF_UI_VERSION - Rename webui-download.cmake -> ui-download.cmake (internal refs updated) - Update labeler.yml: server/webui -> server/ui path label * docs: update CODEOWNERS and server README docs - Update CODEOWNERS: team ggml-org/llama-webui -> ggml-org/llama-ui, path /tools/server/webui/ -> /tools/ui/ - Update server README.md: CLI tables show --ui flags with deprecated --webui aliases - Update server README-dev.md: "WebUI" -> "UI", paths updated to tools/ui/ * fix: Small fixes for UI build * fix: CMake.txt syntax * chore: Formatting * fix: `.editorconfig` for llama-ui * chore: Formatting * refactor: Use `APP_NAME` in Error route * refactor: Cleanup * refactor: Single migration service * make llama-ui a linkable target * fix: UI Build output * fix: Missing change * fix: separate llama-ui npm build output into build/tools/ui/dist subfolder + use cmake npm build instead of downloading ui-build.yml artifacts in CI * refactor: UI workflows cleanup --------- Co-authored-by: Xuan Son Nguyen <son@huggingface.co>
228 lines
6.4 KiB
TypeScript
228 lines
6.4 KiB
TypeScript
import { AgenticSectionType, MessageRole } from '$lib/enums';
|
|
import { ATTACHMENT_SAVED_REGEX, NEWLINE_SEPARATOR } from '$lib/constants';
|
|
import type { ApiChatCompletionToolCall } from '$lib/types/api';
|
|
import type {
|
|
DatabaseMessage,
|
|
DatabaseMessageExtra,
|
|
DatabaseMessageExtraImageFile
|
|
} from '$lib/types/database';
|
|
import { AttachmentType } from '$lib/enums';
|
|
|
|
/**
|
|
* Represents a parsed section of agentic content for display
|
|
*/
|
|
export interface AgenticSection {
|
|
type: AgenticSectionType;
|
|
content: string;
|
|
toolName?: string;
|
|
toolArgs?: string;
|
|
toolResult?: string;
|
|
toolResultExtras?: DatabaseMessageExtra[];
|
|
}
|
|
|
|
/**
|
|
* Represents a tool result line that may reference an image attachment
|
|
*/
|
|
export type ToolResultLine = {
|
|
text: string;
|
|
image?: DatabaseMessageExtraImageFile;
|
|
};
|
|
|
|
/**
|
|
* Derives display sections from a single assistant message and its direct tool results.
|
|
*
|
|
* @param message - The assistant message
|
|
* @param toolMessages - Tool result messages for this assistant's tool_calls
|
|
* @param streamingToolCalls - Partial tool calls during streaming (not yet persisted)
|
|
*/
|
|
function deriveSingleTurnSections(
|
|
message: DatabaseMessage,
|
|
toolMessages: DatabaseMessage[] = [],
|
|
streamingToolCalls: ApiChatCompletionToolCall[] = [],
|
|
isStreaming: boolean = false
|
|
): AgenticSection[] {
|
|
const sections: AgenticSection[] = [];
|
|
|
|
// 1. Reasoning content (from dedicated field)
|
|
if (message.reasoningContent) {
|
|
const toolCalls = parseToolCalls(message.toolCalls);
|
|
const hasContentAfterReasoning =
|
|
!!message.content?.trim() || toolCalls.length > 0 || streamingToolCalls.length > 0;
|
|
const isPending = isStreaming && !hasContentAfterReasoning;
|
|
sections.push({
|
|
type: isPending ? AgenticSectionType.REASONING_PENDING : AgenticSectionType.REASONING,
|
|
content: message.reasoningContent
|
|
});
|
|
}
|
|
|
|
// 2. Text content
|
|
if (message.content?.trim()) {
|
|
sections.push({
|
|
type: AgenticSectionType.TEXT,
|
|
content: message.content
|
|
});
|
|
}
|
|
|
|
// 3. Persisted tool calls (from message.toolCalls field)
|
|
const toolCalls = parseToolCalls(message.toolCalls);
|
|
for (const tc of toolCalls) {
|
|
const resultMsg = toolMessages.find((m) => m.toolCallId === tc.id);
|
|
// Only show as pending/loading if we're actively streaming; otherwise it's just a tool call without result
|
|
const type = resultMsg
|
|
? AgenticSectionType.TOOL_CALL
|
|
: isStreaming
|
|
? AgenticSectionType.TOOL_CALL_PENDING
|
|
: AgenticSectionType.TOOL_CALL;
|
|
sections.push({
|
|
type,
|
|
content: resultMsg?.content || '',
|
|
toolName: tc.function?.name,
|
|
toolArgs: tc.function?.arguments,
|
|
toolResult: resultMsg?.content,
|
|
toolResultExtras: resultMsg?.extra
|
|
});
|
|
}
|
|
|
|
// 4. Streaming tool calls (not yet persisted - currently being received)
|
|
for (const tc of streamingToolCalls) {
|
|
// Skip if already in persisted tool calls
|
|
if (tc.id && toolCalls.find((t) => t.id === tc.id)) continue;
|
|
sections.push({
|
|
type: AgenticSectionType.TOOL_CALL_STREAMING,
|
|
content: '',
|
|
toolName: tc.function?.name,
|
|
toolArgs: tc.function?.arguments
|
|
});
|
|
}
|
|
|
|
return sections;
|
|
}
|
|
|
|
/**
|
|
* Derives display sections from structured message data.
|
|
*
|
|
* Handles both single-turn (one assistant + its tool results) and multi-turn
|
|
* agentic sessions (multiple assistant + tool messages grouped together).
|
|
*
|
|
* When `toolMessages` contains continuation assistant messages (from multi-turn
|
|
* agentic flows), they are processed in order to produce sections across all turns.
|
|
*
|
|
* @param message - The first/anchor assistant message
|
|
* @param toolMessages - Tool result messages and continuation assistant messages
|
|
* @param streamingToolCalls - Partial tool calls during streaming (not yet persisted)
|
|
* @param isStreaming - Whether the message is currently being streamed
|
|
*/
|
|
export function deriveAgenticSections(
|
|
message: DatabaseMessage,
|
|
toolMessages: DatabaseMessage[] = [],
|
|
streamingToolCalls: ApiChatCompletionToolCall[] = [],
|
|
isStreaming: boolean = false
|
|
): AgenticSection[] {
|
|
const hasAssistantContinuations = toolMessages.some((m) => m.role === MessageRole.ASSISTANT);
|
|
|
|
if (!hasAssistantContinuations) {
|
|
return deriveSingleTurnSections(message, toolMessages, streamingToolCalls, isStreaming);
|
|
}
|
|
|
|
const sections: AgenticSection[] = [];
|
|
|
|
const firstTurnToolMsgs = collectToolMessages(toolMessages, 0);
|
|
sections.push(...deriveSingleTurnSections(message, firstTurnToolMsgs));
|
|
|
|
let i = firstTurnToolMsgs.length;
|
|
|
|
while (i < toolMessages.length) {
|
|
const msg = toolMessages[i];
|
|
|
|
if (msg.role === MessageRole.ASSISTANT) {
|
|
const turnToolMsgs = collectToolMessages(toolMessages, i + 1);
|
|
const isLastTurn = i + 1 + turnToolMsgs.length >= toolMessages.length;
|
|
|
|
sections.push(
|
|
...deriveSingleTurnSections(
|
|
msg,
|
|
turnToolMsgs,
|
|
isLastTurn ? streamingToolCalls : [],
|
|
isLastTurn && isStreaming
|
|
)
|
|
);
|
|
|
|
i += 1 + turnToolMsgs.length;
|
|
} else {
|
|
i++;
|
|
}
|
|
}
|
|
|
|
return sections;
|
|
}
|
|
|
|
/**
|
|
* Collect consecutive tool messages starting at `startIndex`.
|
|
*/
|
|
function collectToolMessages(messages: DatabaseMessage[], startIndex: number): DatabaseMessage[] {
|
|
const result: DatabaseMessage[] = [];
|
|
|
|
for (let i = startIndex; i < messages.length; i++) {
|
|
if (messages[i].role === MessageRole.TOOL) {
|
|
result.push(messages[i]);
|
|
} else {
|
|
break;
|
|
}
|
|
}
|
|
|
|
return result;
|
|
}
|
|
|
|
/**
|
|
* Parse tool result text into lines, matching image attachments by name.
|
|
*/
|
|
export function parseToolResultWithImages(
|
|
toolResult: string,
|
|
extras?: DatabaseMessageExtra[]
|
|
): ToolResultLine[] {
|
|
const lines = toolResult.split(NEWLINE_SEPARATOR);
|
|
return lines.map((line) => {
|
|
const match = line.match(ATTACHMENT_SAVED_REGEX);
|
|
if (!match || !extras) return { text: line };
|
|
|
|
const attachmentName = match[1];
|
|
const image = extras.find(
|
|
(e): e is DatabaseMessageExtraImageFile =>
|
|
e.type === AttachmentType.IMAGE && e.name === attachmentName
|
|
);
|
|
|
|
return { text: line, image };
|
|
});
|
|
}
|
|
|
|
/**
|
|
* Safely parse the toolCalls JSON string from a DatabaseMessage.
|
|
*/
|
|
function parseToolCalls(toolCallsJson?: string): ApiChatCompletionToolCall[] {
|
|
if (!toolCallsJson) return [];
|
|
|
|
try {
|
|
const parsed = JSON.parse(toolCallsJson);
|
|
|
|
return Array.isArray(parsed) ? parsed : [];
|
|
} catch {
|
|
return [];
|
|
}
|
|
}
|
|
|
|
/**
|
|
* Check if a message has agentic content (tool calls or is part of an agentic flow).
|
|
*/
|
|
export function hasAgenticContent(
|
|
message: DatabaseMessage,
|
|
toolMessages: DatabaseMessage[] = []
|
|
): boolean {
|
|
if (message.toolCalls) {
|
|
const tc = parseToolCalls(message.toolCalls);
|
|
|
|
if (tc.length > 0) return true;
|
|
}
|
|
|
|
return toolMessages.length > 0;
|
|
}
|