mirror of
https://github.com/DayuanJiang/next-ai-draw-io.git
synced 2026-10-11 12:09:53 +08:00
Found by the second PR review: - With AWS_BEARER_TOKEN_BEDROCK set on the server, a request with the user's AWS keys ran on the server's token: the Bedrock SDK prefers it. Checked with Bedrock: invalid user keys used to get an answer. - An OpenAI key with the official URL filled in (the settings form does that) went to the Responses API. Back to main's rule: a configured base URL uses Chat Completions. - A user's Ollama key went to the server's OLLAMA_BASE_URL, for chat and for the model list. Like every other provider, it goes to the user's base URL or Ollama Cloud. - The server's keyless Ollama and EdgeOne were not counted in the quota. - AI_MODEL models on the server's keys ran on any provider with a server key, not only on AI_PROVIDER. - A user's Azure key without a base URL used the server's resource name. - The admin panel's Test button failed whenever access codes were set. - DeepSeek's errors in the stream (plain text) were shown as they were, without a hint and also on the server's keys. Bedrock's throttling in the stream was not recognised as a rate limit. - The EdgeOne function accepted text/plain; x=application/json, which other sites can send without a CORS preflight. - Desktop app: a launch that found the old port taken for a moment (the previous version still quitting after an update) remembered the new port for good. The new port is kept only when Windows reserves the old one. A failed read of the presets file moved it aside as corrupt, and a save could then replace the presets. Switching presets on the same port now reloads the page. The dev launcher no longer misses a preset change made before or during a restart.
140 lines
4.1 KiB
TypeScript
140 lines
4.1 KiB
TypeScript
// @vitest-environment node
|
|
import { afterEach, beforeEach, describe, expect, it, vi } from "vitest"
|
|
|
|
// Quota on, and every check answers that the daily limit is used up
|
|
const quota = vi.hoisted(() => ({ checks: 0 }))
|
|
vi.mock("@/lib/dynamo-quota-manager", () => ({
|
|
isQuotaEnabled: () => true,
|
|
checkAndIncrementRequest: async () => {
|
|
quota.checks++
|
|
return {
|
|
allowed: false,
|
|
error: "Daily limit reached",
|
|
type: "request",
|
|
used: 10,
|
|
limit: 10,
|
|
}
|
|
},
|
|
recordTokenUsage: async () => {},
|
|
}))
|
|
|
|
import { POST as chat } from "@/app/api/chat/route"
|
|
|
|
const ENV = [
|
|
"AI_PROVIDER",
|
|
"AI_MODEL",
|
|
"OPENAI_API_KEY",
|
|
"OLLAMA_BASE_URL",
|
|
"OLLAMA_API_KEY",
|
|
"AI_GATEWAY_API_KEY",
|
|
"ALLOW_PRIVATE_URLS",
|
|
]
|
|
const saved: Record<string, string | undefined> = {}
|
|
|
|
beforeEach(() => {
|
|
for (const k of ENV) saved[k] = process.env[k]
|
|
process.env.AI_PROVIDER = "openai"
|
|
process.env.AI_MODEL = "gpt-5.5"
|
|
process.env.OPENAI_API_KEY = "server-key"
|
|
quota.checks = 0
|
|
// No request may reach a provider
|
|
vi.stubGlobal(
|
|
"fetch",
|
|
vi.fn(async () => {
|
|
throw new Error("no network in tests")
|
|
}),
|
|
)
|
|
})
|
|
|
|
afterEach(() => {
|
|
for (const k of ENV) {
|
|
if (saved[k] === undefined) delete process.env[k]
|
|
else process.env[k] = saved[k]
|
|
}
|
|
vi.unstubAllGlobals()
|
|
})
|
|
|
|
const send = (headers: Record<string, string>) =>
|
|
chat(
|
|
new Request("http://localhost/api/chat", {
|
|
method: "POST",
|
|
headers: {
|
|
"Content-Type": "application/json",
|
|
"x-forwarded-for": "203.0.113.7",
|
|
...headers,
|
|
},
|
|
body: JSON.stringify({
|
|
messages: [
|
|
{
|
|
id: "u1",
|
|
role: "user",
|
|
parts: [{ type: "text", text: "Draw two boxes" }],
|
|
},
|
|
],
|
|
xml: "",
|
|
}),
|
|
}),
|
|
)
|
|
|
|
describe("chat quota", () => {
|
|
it("counts a request whose key header the provider never reads", async () => {
|
|
// OpenAI ignores the AWS key, so this runs on the server's key
|
|
const res = await send({
|
|
"x-ai-provider": "openai",
|
|
"x-aws-access-key-id": "x",
|
|
})
|
|
expect(res.status).toBe(429)
|
|
expect(quota.checks).toBe(1)
|
|
})
|
|
|
|
it("does not count a request on the user's own key", async () => {
|
|
const res = await send({
|
|
"x-ai-provider": "openai",
|
|
"x-ai-api-key": "user-key",
|
|
"x-ai-model": "gpt-5.5",
|
|
})
|
|
expect(res.status).not.toBe(429)
|
|
expect(quota.checks).toBe(0)
|
|
})
|
|
|
|
it("counts the server's keyless Ollama and EdgeOne", async () => {
|
|
process.env.AI_PROVIDER = "ollama"
|
|
process.env.AI_MODEL = "llama3.2"
|
|
process.env.OLLAMA_BASE_URL = "http://ollama.internal:11434/api"
|
|
expect((await send({})).status).toBe(429)
|
|
expect(
|
|
(
|
|
await send({
|
|
"x-ai-provider": "edgeone",
|
|
"x-ai-model": "@tx/deepseek-ai/deepseek-v3-0324",
|
|
})
|
|
).status,
|
|
).toBe(429)
|
|
expect(quota.checks).toBe(2)
|
|
})
|
|
|
|
it("does not count Ollama on the user's own server", async () => {
|
|
const res = await send({
|
|
"x-ai-provider": "ollama",
|
|
"x-ai-base-url": "https://ollama.example.com/api",
|
|
"x-ai-model": "llama3.2",
|
|
})
|
|
expect(res.status).not.toBe(429)
|
|
expect(quota.checks).toBe(0)
|
|
})
|
|
})
|
|
|
|
describe("server model allowlist", () => {
|
|
it("runs AI_MODEL only on the server's AI_PROVIDER", async () => {
|
|
// Another provider's server key must not run it
|
|
process.env.AI_GATEWAY_API_KEY = "server-gateway-key"
|
|
const res = await send({
|
|
"x-ai-provider": "gateway",
|
|
"x-ai-model": "gpt-5.5",
|
|
})
|
|
expect(res.status).toBe(400)
|
|
expect(await res.text()).toMatch(/not available on this server/)
|
|
expect(quota.checks).toBe(0)
|
|
})
|
|
})
|