mirror of
https://github.com/diegosouzapw/OmniRoute.git
synced 2026-08-07 07:42:13 +03:00
* chore(release): open v3.8.22 development cycle * refactor(dashboard): extract ProviderDetailPageClient — #3501 Phase 0 (#3633) #3501 Phase 0: extract ProviderDetailPageClient + smoke test. Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * refactor(dashboard): extract auth-import modals — #3501 Phase 1a (#3634) #3501 Phase 1a: extract 3 auth-import modal clusters. Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * fix(db): reclassify localDb unexported modules as intentionally-internal (#3499) (#3635) Closes #3499 — reclassify localDb unexported modules as intentionally-internal (audit + honest gate framing). * refactor(db): move call_logs aggregations into callLogStats db module (#3500) (#3636) #3500 slice 1: call_logs aggregations → src/lib/db/callLogStats.ts (Rule #5). Byte-identical queries; TDD 6/6. * refactor(dashboard): extract EditCompatibleNodeModal — #3501 Phase 1b (#3638) #3501 Phase 1b: extract EditCompatibleNodeModal (cycle-safe via leaf constants module). Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * refactor(db): move community_servers SQL into gamification db module (#3500 slice 3) (#3639) #3500 slice 3: community_servers SQL → gamification db module. * refactor(db): move usage_history SQL into usageAnalytics module (#3500 slice 2) (#3644) #3500 slice 2: usage_history/daily_usage_summary SQL → usageAnalytics db module. * refactor(db): move skills UPDATE + db-backups SQL into db modules (#3500 slice 5) (#3647) #3500 slice 5: skills UPDATE (allowlist) + db-backups SQL → db modules. * refactor(db): move usage_logs/semantic_cache/proxy_logs SQL into db modules (#3500 slice 4) (#3648) #3500 slice 4: usage_logs/semantic_cache/proxy_logs SQL → db modules. All internal routes done (2 external by-design remain). * chore(db-gate): reclassify external-DB reads, fully close #3500 (#3649) Closes #3500: reclassify external-DB reads; all internal raw-SQL migrated to db/ modules. * refactor(dashboard): extract pure helpers to providerPageHelpers — #3501 Phase 2 (#3653) #3501 Phase 2: extract pure helpers to providerPageHelpers (leaf, cycle-safe). Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * refactor(dashboard): extract remaining shared helpers to providerPageHelpers — #3501 Phase 2b (#3658) #3501 Phase 2b: extract remaining shared helpers to providerPageHelpers (leaf, cycle-safe). Heavy modals unblocked. Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * fix(reasoning): replay reasoning_content on plain DeepSeek turns (#1682) (#3632) Integrated into release/v3.8.22 * fix(kiro): route enterprise IAM Identity Center accounts to their regional endpoint (#3631) Integrated into release/v3.8.22 * refactor: small code cleanup (#3523) Integrated into release/v3.8.22 * fix(combo): skip same-provider targets on 408/500/502/503/504/524 errors (#3637) Integrated into release/v3.8.22 — circuit-breaker guard added in review (#1731v2) * feat(providers): add MiMoCode free-tier provider with bootstrap JWT auth (#3659) Integrated into release/v3.8.22 — page.tsx conflict resolved + NoAuthAccountCard re-applied to ProviderDetailPageClient in review. MiMoCode endpoint validated live. * Log Responses WebSocket calls in history (#3616) Integrated into release/v3.8.22 — Codex Responses WebSocket call history logging. * Add Claude Code routing preference for unprefixed Claude models (#3540) Integrated into release/v3.8.22 — page.tsx conflict resolved (re-applied toggle to ProviderDetailPageClient) + disable-test updated for catalog drift in review. * docs(changelog): credit #3632/#3631/#3637/#3659/#3540/#3616/#3523 (v3.8.22 targeted review round) * fix(mimocode): add required authHeader:"none" to registry entry (#3659 follow-up) The mimocode RegistryEntry omitted the required authHeader field, which broke typecheck:core (TS2741). Match the no-auth convention (authType:"none" + authHeader:"none") used by veoaifree-web and other free providers. Follow-up to #3659 (@pizzav-xyz). * fix(responses): detect stream readiness for tool-call-only and object-less chunks (#3612) (#3661) Closes #3612 * fix(mitm): remove duplicated 'Command failed:' error prefix (#3641) (#3662) Closes #3641 * fix(cli): honor HERMES_HOME for Hermes Agent config path (#3628) (#3663) Closes #3628 * fix(api): fetch live OpenCode model catalog for no-auth model picker (#3611) (#3664) Closes #3611 * fix(api): flag provider topology error state by current status, not stale history (#3619) (#3666) Closes #3619 * fix(electron): launch peer-stamping server-ws.mjs entrypoint to avoid 403 LOCAL_ONLY (#3386) (#3665) Closes #3386 * fix(dashboard): restore home topology live in-flight pulse (#3507) (#3667) Closes #3507 * fix(oauth): name Kiro/AWS auto-imported accounts and dedupe by profileArn (#3615) (#3671) Closes #3615 * fix(resilience): clear stale transient connection cooldowns on startup (#3625) (#3672) Closes #3625 * fix(i18n): use logical CSS direction utilities for sidebar and key overlays (RTL #3541) (#3670) Closes #3541 * fix(dashboard): honor auto-hide and switch to visible filter on passthrough Test-all (#3610) (#3669) Closes #3610 * refactor(dashboard): extract AddApiKeyModal + EditConnectionModal — #3501 Phase 1c (#3674) #3501 Phase 1c: extract AddApiKeyModal, EditConnectionModal, WebSessionCredentialGuide into components/; god-component 10,166->8,092 LOC. Reconciles the v3.8.22 file-size drift for this file. Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * docs(changelog): reconcile v3.8.22 — credit #3621/#3622 + MiMoCode follow-up roll-up * refactor(dashboard): extract ConnectionRow + ModelCompatPopover + SiliconFlowEndpointModal — #3501 Phase 1d (#3676) #3501 Phase 1d: god-component 8,092->6,838 LOC. Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * feat(obsidian): add WebDAV config route + encrypt creds at rest (#3485 part 1) (#3677) Part 1 of #3485. Adds /api/settings/obsidian/webdav (GET/POST/DELETE) wiring the ready obsidianSync lib, encrypts webdav password + obsidian token at rest, removes the duplicate UI block, drops the KNOWN_MISSING entry. WebDAV file server is part 2. * feat(obsidian): add /api/v1/webdav file server for Obsidian vault sync (#3485 part 2) (#3678) Part 2 of #3485. WebDAV server (PROPFIND/GET/PUT/DELETE/MKCOL/MOVE/OPTIONS) handled in the custom server layer (standalone-server-ws.mjs) since the App Router cannot export WebDAV methods. Basic-Auth (constant-time), path-traversal hardened, password decrypt ported from encryption.ts (parity-tested), DATA_DIR resolution parity-tested against dataPaths.ts. End-to-end Obsidian-over-Tailscale validation is a live VPS step (Rule #18). * fix(combo): stop premature context compaction — real auto-combo windows + per-target compression limit (#3680) Integrated into release/v3.8.22 * feat(dashboard): deactivate/activate accounts from the quota overview (#3675) Integrated into release/v3.8.22 * fix(dashboard): close review gaps in bulk provider connection actions (#3271 follow-up) (#3673) Integrated into release/v3.8.22 — page.tsx conflict (god-component split #3501) resolved by re-applying the bulk-action deltas to ProviderDetailPageClient.tsx * refactor(dashboard): extract useModelCompatState hook + model sections — #3501 Phase 1e (#3683) #3501 Phase 1e: extract useModelCompatState hook (unblocks the model sections) + ModelRow/PassthroughModelsSection/PassthroughModelRow/CustomModelsSection/CompatibleModelsSection. god-component 6,838->4,921 LOC. Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * refactor(dashboard): extract useProviderConnections/Settings/Models hooks — #3501 Phase 1f (#3684) #3501 Phase 1f: god-component 4,948->4,062 LOC. Connection state+handlers, settings, and model metadata moved into hooks/. Co-authored-by: oyi77 <oyi77@users.noreply.github.com> * chore(release): v3.8.22 CHANGELOG + env-doc sync - Set release date in CHANGELOG [3.8.22] to 2026-06-11 - Add HERMES_HOME to .env.example (from #3628/#3663) - Add HERMES_HOME + OMNIROUTE_PREFER_CLAUDE_CODE_FOR_UNPREFIXED_CLAUDE_MODELS to ENVIRONMENT.md (#3628/#3540) * docs(changelog): credit #3673 + #3675 — leninejunior bulk-actions + quota-toggle --------- Co-authored-by: oyi77 <oyi77@users.noreply.github.com> Co-authored-by: Abhishek Divekar <adivekar@utexas.edu> Co-authored-by: NOXX - Commiter <artur1992123@mail.ru> Co-authored-by: Nicolas Lorin <androw95220@gmail.com> Co-authored-by: Hernan Javier Ardila Sanchez <hjasgr@gmail.com> Co-authored-by: PizzaV <103120356+pizzav-xyz@users.noreply.github.com> Co-authored-by: kkkayye <98376609+kkkayye@users.noreply.github.com> Co-authored-by: Witroch4 <witalo_rocha@hotmail.com> Co-authored-by: Lenine Júnior <lenine@engrene.com.br>
329 lines
11 KiB
TypeScript
329 lines
11 KiB
TypeScript
import { AutoComboConfig } from "./engine";
|
|
import { MODE_PACKS } from "./modePacks";
|
|
import { DEFAULT_WEIGHTS, ScoringWeights } from "./scoring";
|
|
import { AutoVariant } from "./autoPrefix";
|
|
import { getProviderConnections } from "@/lib/db/providers";
|
|
import { getProviderRegistry } from "./providerRegistryAccessor";
|
|
import type { ConnectionFields } from "@/lib/db/encryption";
|
|
import { NOAUTH_PROVIDERS } from "@/shared/constants/providers";
|
|
import { hasUsableWebSessionCredential } from "@/shared/providers/webSessionCredentials";
|
|
import { defaultLogger as log } from "@omniroute/open-sse/utils/logger";
|
|
import { getTokenLimit } from "../contextManager";
|
|
import { getResolvedModelCapabilities } from "@/lib/modelCapabilities";
|
|
|
|
/** Minimal connection shape needed for virtual auto-combo factory */
|
|
interface VirtualFactoryConn extends ConnectionFields {
|
|
id: string;
|
|
provider: string;
|
|
defaultModel?: string;
|
|
expiresAt?: number | string | null;
|
|
tokenExpiresAt?: number | string | null;
|
|
providerSpecificData?: Record<string, unknown> | null;
|
|
}
|
|
|
|
type NoAuthProviderDefinition = {
|
|
id?: string;
|
|
alias?: string;
|
|
noAuth?: boolean;
|
|
serviceKinds?: string[];
|
|
};
|
|
|
|
export interface VirtualAutoComboCandidate {
|
|
provider: string;
|
|
connectionId: string;
|
|
model: string;
|
|
modelStr: string; // e.g., 'openai/gpt-4o'
|
|
costPer1MTokens: number; // from providerRegistry
|
|
}
|
|
|
|
type VirtualAutoCombo = AutoComboConfig & {
|
|
strategy: "auto";
|
|
models: Array<{
|
|
id: string;
|
|
kind: "model";
|
|
model: string;
|
|
providerId: string;
|
|
connectionId: string;
|
|
weight: number;
|
|
label: string;
|
|
}>;
|
|
/** MAX of candidates' context windows — safe to advertise because the
|
|
* auto-combo context pre-filter routes oversized requests to large-window
|
|
* candidates. null when the pool is empty. */
|
|
advertisedContextLength: number | null;
|
|
advertisedMaxOutputTokens: number | null;
|
|
autoConfig: {
|
|
candidatePool: string[];
|
|
weights: ScoringWeights;
|
|
explorationRate: number;
|
|
routerStrategy: string;
|
|
};
|
|
config: {
|
|
auto: {
|
|
candidatePool: string[];
|
|
weights: ScoringWeights;
|
|
explorationRate: number;
|
|
routerStrategy: string;
|
|
};
|
|
};
|
|
};
|
|
|
|
function toExpiryMs(value: unknown): number | null {
|
|
if (value === null || value === undefined || value === "") return null;
|
|
|
|
const parsed =
|
|
typeof value === "number"
|
|
? value
|
|
: typeof value === "string" && value.trim().length > 0
|
|
? Number(value)
|
|
: Number.NaN;
|
|
|
|
if (Number.isFinite(parsed) && parsed > 0) {
|
|
return parsed < 10_000_000_000 ? parsed * 1000 : parsed;
|
|
}
|
|
|
|
if (typeof value === "string") {
|
|
const timestamp = new Date(value).getTime();
|
|
return Number.isFinite(timestamp) ? timestamp : null;
|
|
}
|
|
|
|
return null;
|
|
}
|
|
|
|
function hasUsableOAuthToken(conn: VirtualFactoryConn): boolean {
|
|
if (typeof conn.accessToken !== "string" || conn.accessToken.trim().length === 0) return false;
|
|
|
|
const expiryMs = toExpiryMs(conn.tokenExpiresAt) ?? toExpiryMs(conn.expiresAt);
|
|
|
|
return expiryMs === null || expiryMs > Date.now();
|
|
}
|
|
|
|
function hasProviderSpecificSessionData(conn: VirtualFactoryConn): boolean {
|
|
return hasUsableWebSessionCredential(conn.provider, conn.providerSpecificData);
|
|
}
|
|
|
|
function hasUsableConnectionCredential(conn: VirtualFactoryConn): boolean {
|
|
const hasApiKey = typeof conn.apiKey === "string" && conn.apiKey.trim().length > 0;
|
|
return hasApiKey || hasUsableOAuthToken(conn) || hasProviderSpecificSessionData(conn);
|
|
}
|
|
|
|
const SYNTHETIC_NOAUTH_CONNECTION_ID = "noauth";
|
|
|
|
function isChatAutoComboNoAuthProvider(providerDef: NoAuthProviderDefinition): boolean {
|
|
if (providerDef.noAuth !== true) return false;
|
|
if (!Array.isArray(providerDef.serviceKinds) || providerDef.serviceKinds.length === 0)
|
|
return true;
|
|
return providerDef.serviceKinds.includes("llm");
|
|
}
|
|
|
|
function getFirstRegistryModelId(providerInfo: { models?: Array<{ id?: string }> } | undefined) {
|
|
const firstModel = Array.isArray(providerInfo?.models) ? providerInfo.models[0] : undefined;
|
|
return typeof firstModel?.id === "string" && firstModel.id.trim().length > 0
|
|
? firstModel.id
|
|
: undefined;
|
|
}
|
|
|
|
function getNoAuthCandidates(excludedProviders: Set<string>): VirtualAutoComboCandidate[] {
|
|
const registry = getProviderRegistry();
|
|
const candidates: VirtualAutoComboCandidate[] = [];
|
|
|
|
for (const providerDef of Object.values(NOAUTH_PROVIDERS) as NoAuthProviderDefinition[]) {
|
|
if (!isChatAutoComboNoAuthProvider(providerDef)) continue;
|
|
|
|
const providerId = providerDef.id;
|
|
if (!providerId || excludedProviders.has(providerId)) continue;
|
|
|
|
const providerInfo = registry[providerId];
|
|
const modelId = getFirstRegistryModelId(providerInfo);
|
|
if (!modelId) continue;
|
|
|
|
// No-auth providers do not have provider_connections rows. Use the same
|
|
// synthetic connection id returned by getProviderCredentials() so the
|
|
// downstream combo path can still carry a stable target/account identity.
|
|
// Prefer provider aliases because some canonical provider IDs are reserved
|
|
// for credentialed tiers with different routing semantics.
|
|
const registryAlias =
|
|
typeof providerInfo?.alias === "string" && providerInfo.alias.trim().length > 0
|
|
? providerInfo.alias
|
|
: null;
|
|
const routingPrefix = providerDef.alias || registryAlias || providerId;
|
|
candidates.push({
|
|
provider: providerId,
|
|
connectionId: SYNTHETIC_NOAUTH_CONNECTION_ID,
|
|
model: modelId,
|
|
modelStr: `${routingPrefix}/${modelId}`,
|
|
costPer1MTokens: 0,
|
|
});
|
|
}
|
|
|
|
return candidates;
|
|
}
|
|
|
|
/**
|
|
* Creates a virtual AutoCombo configuration dynamically based on connected providers and a specified variant.
|
|
* This combo is not persisted in the DB.
|
|
*/
|
|
/**
|
|
* Aggregate the context window / max output to ADVERTISE for an auto combo.
|
|
*
|
|
* MAX across candidates (not min): the auto-combo context pre-filter
|
|
* (combo.ts::filterTargetsByRequestCompatibility + the estimated-tokens
|
|
* pre-filter) already routes oversized requests away from small-window
|
|
* candidates, so advertising the largest window lets clients (e.g. opencode)
|
|
* keep their smart auto-compaction calibrated to the best candidate instead
|
|
* of compacting prematurely — or, worse, receiving 0 and disabling
|
|
* compaction entirely (the "agent keeps forgetting things" bug).
|
|
*
|
|
* Unknown candidates resolve through getTokenLimit()'s fallback chain, so a
|
|
* non-empty pool always yields a positive contextLength.
|
|
*/
|
|
export function computeAdvertisedLimits(candidates: Array<{ provider: string; model: string }>): {
|
|
contextLength: number | null;
|
|
maxOutputTokens: number | null;
|
|
} {
|
|
if (!Array.isArray(candidates) || candidates.length === 0) {
|
|
return { contextLength: null, maxOutputTokens: null };
|
|
}
|
|
|
|
let contextLength: number | null = null;
|
|
let maxOutputTokens: number | null = null;
|
|
for (const candidate of candidates) {
|
|
const limit = getTokenLimit(candidate.provider, candidate.model);
|
|
if (Number.isFinite(limit) && limit > 0) {
|
|
contextLength = contextLength === null ? limit : Math.max(contextLength, limit);
|
|
}
|
|
const output = getResolvedModelCapabilities({
|
|
provider: candidate.provider,
|
|
model: candidate.model,
|
|
}).maxOutputTokens;
|
|
if (typeof output === "number" && Number.isFinite(output) && output > 0) {
|
|
maxOutputTokens = maxOutputTokens === null ? output : Math.max(maxOutputTokens, output);
|
|
}
|
|
}
|
|
return { contextLength, maxOutputTokens };
|
|
}
|
|
|
|
export async function createVirtualAutoCombo(
|
|
variant: AutoVariant | undefined
|
|
): Promise<VirtualAutoCombo> {
|
|
const connections = (await getProviderConnections({ isActive: true })) as VirtualFactoryConn[];
|
|
|
|
const validConnections = connections.filter(hasUsableConnectionCredential);
|
|
|
|
const candidatePool: VirtualAutoComboCandidate[] = [];
|
|
for (const conn of validConnections) {
|
|
const providerInfo = getProviderRegistry()[conn.provider];
|
|
if (!providerInfo) continue; // Skip unknown providers
|
|
|
|
let modelId: string | undefined = conn.defaultModel;
|
|
if (!modelId) {
|
|
const firstModel = providerInfo.models[0];
|
|
modelId = firstModel?.id;
|
|
}
|
|
if (!modelId) continue; // Skip providers without a model
|
|
|
|
candidatePool.push({
|
|
provider: conn.provider,
|
|
connectionId: conn.id,
|
|
model: modelId,
|
|
modelStr: `${conn.provider}/${modelId}`,
|
|
costPer1MTokens: 0, // Not used in virtual auto-combo (LKGP uses session stickiness)
|
|
});
|
|
}
|
|
|
|
candidatePool.push(
|
|
...getNoAuthCandidates(new Set(validConnections.map((conn) => conn.provider)))
|
|
);
|
|
|
|
if (candidatePool.length === 0) {
|
|
log.warn("AUTO", "No connected providers with valid credentials for virtual auto-combo");
|
|
const emptyPool: string[] = [];
|
|
const autoConfig = {
|
|
candidatePool: emptyPool,
|
|
weights: { ...DEFAULT_WEIGHTS },
|
|
explorationRate: 0.05,
|
|
routerStrategy: "lkgp",
|
|
};
|
|
return {
|
|
id: `virtual-auto-${variant || "default"}`,
|
|
name: `Auto ${variant || "Default"}`,
|
|
type: "auto" as const,
|
|
strategy: "auto",
|
|
models: [],
|
|
candidatePool: emptyPool,
|
|
weights: autoConfig.weights,
|
|
explorationRate: autoConfig.explorationRate,
|
|
routerStrategy: autoConfig.routerStrategy,
|
|
autoConfig,
|
|
config: { auto: autoConfig },
|
|
advertisedContextLength: null,
|
|
advertisedMaxOutputTokens: null,
|
|
};
|
|
}
|
|
|
|
let weights: ScoringWeights = { ...DEFAULT_WEIGHTS };
|
|
let explorationRate = 0.05; // Default exploration rate
|
|
let routerStrategy = "lkgp"; // All auto variants use LKGP
|
|
|
|
switch (variant) {
|
|
case "coding":
|
|
weights = { ...MODE_PACKS["quality-first"] };
|
|
break;
|
|
case "fast":
|
|
weights = { ...MODE_PACKS["ship-fast"] };
|
|
break;
|
|
case "cheap":
|
|
weights = { ...MODE_PACKS["cost-saver"] };
|
|
break;
|
|
case "offline":
|
|
weights = { ...MODE_PACKS["offline-friendly"] };
|
|
break;
|
|
case "smart":
|
|
weights = { ...MODE_PACKS["quality-first"] };
|
|
explorationRate = 0.1; // Override default exploration rate
|
|
break;
|
|
case "lkgp":
|
|
// LKGP is default for all auto variants, this variant just explicitly names it.
|
|
// Use default weights.
|
|
break;
|
|
case undefined: // Default auto
|
|
// Use default weights
|
|
break;
|
|
}
|
|
|
|
const providerPool = [...new Set(candidatePool.map((c) => c.provider))];
|
|
const models = candidatePool.map((candidate, index) => ({
|
|
id: `virtual-auto-${variant || "default"}-${index + 1}-${candidate.provider}`,
|
|
kind: "model" as const,
|
|
model: candidate.modelStr,
|
|
providerId: candidate.provider,
|
|
connectionId: candidate.connectionId,
|
|
weight: 1,
|
|
label: candidate.provider,
|
|
}));
|
|
const autoConfig = {
|
|
candidatePool: providerPool,
|
|
weights,
|
|
explorationRate,
|
|
routerStrategy,
|
|
};
|
|
|
|
const advertisedLimits = computeAdvertisedLimits(candidatePool);
|
|
|
|
return {
|
|
id: `virtual-auto-${variant || "default"}`,
|
|
name: `Auto ${variant || "Default"}`,
|
|
type: "auto",
|
|
strategy: "auto",
|
|
models,
|
|
candidatePool: providerPool,
|
|
weights,
|
|
explorationRate,
|
|
routerStrategy,
|
|
autoConfig,
|
|
config: { auto: autoConfig },
|
|
advertisedContextLength: advertisedLimits.contextLength,
|
|
advertisedMaxOutputTokens: advertisedLimits.maxOutputTokens,
|
|
};
|
|
}
|