NextChat-U/app/client/platforms/openai.ts

"use client";
// azure and openai, using same models. so using same LLMApi.
import {
  ApiPath,
  OPENAI_BASE_URL,
  DEFAULT_MODELS,
  OpenaiPath,
  Azure,
  REQUEST_TIMEOUT_MS,
  ServiceProvider,
} from "@/app/constant";
import {
  ChatMessageTool,
  useAccessStore,
  useAppConfig,
  useChatStore,
  usePluginStore,
} from "@/app/store";
import { collectModelsWithDefaultModel } from "@/app/utils/model";
import {
  preProcessImageContent,
  uploadImage,
  base64Image2Blob,
  stream,
} from "@/app/utils/chat";
import { cloudflareAIGatewayUrl } from "@/app/utils/cloudflare";
import { DalleSize, DalleQuality, DalleStyle } from "@/app/typing";

import {
  ChatOptions,
  getHeaders,
  LLMApi,
  LLMModel,
  LLMUsage,
  MultimodalContent,
  SpeechOptions,
} from "../api";
import Locale from "../../locales";
import { getClientConfig } from "@/app/config/client";
import {
  getMessageTextContent,
  isVisionModel,
  isDalle3 as _isDalle3,
} from "@/app/utils";
import { fetch } from "@/app/utils/stream";

export interface OpenAIListModelResponse {
  object: string;
  data: Array<{
    id: string;
    object: string;
    root: string;
  }>;
}

export interface RequestPayload {
  messages: {
    role: "system" | "user" | "assistant";
    content: string | MultimodalContent[];
  }[];
  stream?: boolean;
  model: string;
  temperature: number;
  presence_penalty: number;
  frequency_penalty: number;
  top_p: number;
  max_tokens?: number;
  max_completion_tokens?: number;
}

export interface DalleRequestPayload {
  model: string;
  prompt: string;
  response_format: "url" | "b64_json";
  n: number;
  size: DalleSize;
  quality: DalleQuality;
  style: DalleStyle;
}

export class ChatGPTApi implements LLMApi {
  private disableListModels = true;

  path(path: string): string {
    const accessStore = useAccessStore.getState();

    let baseUrl = "";

    const isAzure = path.includes("deployments");
    if (accessStore.useCustomConfig) {
      if (isAzure && !accessStore.isValidAzure()) {
        throw Error(
          "incomplete azure config, please check it in your settings page",
        );
      }

      baseUrl = isAzure ? accessStore.azureUrl : accessStore.openaiUrl;
    }

    if (baseUrl.length === 0) {
      const isApp = !!getClientConfig()?.isApp;
      const apiPath = isAzure ? ApiPath.Azure : ApiPath.OpenAI;
      baseUrl = isApp ? OPENAI_BASE_URL : apiPath;
    }

    if (baseUrl.endsWith("/")) {
      baseUrl = baseUrl.slice(0, baseUrl.length - 1);
    }
    if (
      !baseUrl.startsWith("http") &&
      !isAzure &&
      !baseUrl.startsWith(ApiPath.OpenAI)
    ) {
      baseUrl = "https://" + baseUrl;
    }

    console.log("[Proxy Endpoint] ", baseUrl, path);

    // try rebuild url, when using cloudflare ai gateway in client
    return cloudflareAIGatewayUrl([baseUrl, path].join("/"));
  }

  async extractMessage(res: any) {
    if (res.error) {
      return "```\n" + JSON.stringify(res, null, 4) + "\n```";
    }
    // dalle3 model return url, using url create image message
    if (res.data) {
      let url = res.data?.at(0)?.url ?? "";
      const b64_json = res.data?.at(0)?.b64_json ?? "";
      if (!url && b64_json) {
        // uploadImage
        url = await uploadImage(base64Image2Blob(b64_json, "image/png"));
      }
      return [
        {
          type: "image_url",
          image_url: {
            url,
          },
        },
      ];
    }
    return res.choices?.at(0)?.message?.content ?? res;
  }

  async speech(options: SpeechOptions): Promise<ArrayBuffer> {
    const requestPayload = {
      model: options.model,
      input: options.input,
      voice: options.voice,
      response_format: options.response_format,
      speed: options.speed,
    };

    console.log("[Request] openai speech payload: ", requestPayload);

    const controller = new AbortController();
    options.onController?.(controller);

    try {
      const speechPath = this.path(OpenaiPath.SpeechPath);
      const speechPayload = {
        method: "POST",
        body: JSON.stringify(requestPayload),
        signal: controller.signal,
        headers: getHeaders(),
      };

      // make a fetch request
      const requestTimeoutId = setTimeout(
        () => controller.abort(),
        REQUEST_TIMEOUT_MS,
      );

      const res = await fetch(speechPath, speechPayload);
      clearTimeout(requestTimeoutId);
      return await res.arrayBuffer();
    } catch (e) {
      console.log("[Request] failed to make a speech request", e);
      throw e;
    }
  }

  async chat(options: ChatOptions) {
    const modelConfig = {
      ...useAppConfig.getState().modelConfig,
      ...useChatStore.getState().currentSession().mask.modelConfig,
      ...{
        model: options.config.model,
        providerName: options.config.providerName,
      },
    };

    let requestPayload: RequestPayload | DalleRequestPayload;

    const isDalle3 = _isDalle3(options.config.model);
    const isO1 = options.config.model.startsWith("o1");
    if (isDalle3) {
      const prompt = getMessageTextContent(
        options.messages.slice(-1)?.pop() as any,
      );
      requestPayload = {
        model: options.config.model,
        prompt,
        // URLs are only valid for 60 minutes after the image has been generated.
        response_format: "b64_json", // using b64_json, and save image in CacheStorage
        n: 1,
        size: options.config?.size ?? "1024x1024",
        quality: options.config?.quality ?? "standard",
        style: options.config?.style ?? "vivid",
      };
    } else {
      const visionModel = isVisionModel(options.config.model);
      const messages: ChatOptions["messages"] = [];
      for (const v of options.messages) {
        const content = visionModel
          ? await preProcessImageContent(v.content)
          : getMessageTextContent(v);
        if (!(isO1 && v.role === "system"))
          messages.push({ role: v.role, content });
      }

      // O1 not support image, tools (plugin in ChatGPTNextWeb) and system, stream, logprobs, temperature, top_p, n, presence_penalty, frequency_penalty yet.
      requestPayload = {
        messages,
        stream: options.config.stream,
        model: modelConfig.model,
        temperature: !isO1 ? modelConfig.temperature : 1,
        presence_penalty: !isO1 ? modelConfig.presence_penalty : 0,
        frequency_penalty: !isO1 ? modelConfig.frequency_penalty : 0,
        top_p: !isO1 ? modelConfig.top_p : 1,
        // max_tokens: Math.max(modelConfig.max_tokens, 1024),
        // Please do not ask me why not send max_tokens, no reason, this param is just shit, I dont want to explain anymore.
      };

      // O1 使用 max_completion_tokens 控制token数 (https://platform.openai.com/docs/guides/reasoning#controlling-costs)
      if (isO1) {
        requestPayload["max_completion_tokens"] = modelConfig.max_tokens;
      }

      // add max_tokens to vision model
      if (visionModel) {
        requestPayload["max_tokens"] = Math.max(modelConfig.max_tokens, 4000);
      }
    }

    console.log("[Request] openai payload: ", requestPayload);

    const shouldStream = !isDalle3 && !!options.config.stream;
    const controller = new AbortController();
    options.onController?.(controller);

    try {
      let chatPath = "";
      if (modelConfig.providerName === ServiceProvider.Azure) {
        // find model, and get displayName as deployName
        const { models: configModels, customModels: configCustomModels } =
          useAppConfig.getState();
        const {
          defaultModel,
          customModels: accessCustomModels,
          useCustomConfig,
        } = useAccessStore.getState();
        const models = collectModelsWithDefaultModel(
          configModels,
          [configCustomModels, accessCustomModels].join(","),
          defaultModel,
        );
        const model = models.find(
          (model) =>
            model.name === modelConfig.model &&
            model?.provider?.providerName === ServiceProvider.Azure,
        );
        chatPath = this.path(
          (isDalle3 ? Azure.ImagePath : Azure.ChatPath)(
            (model?.displayName ?? model?.name) as string,
            useCustomConfig ? useAccessStore.getState().azureApiVersion : "",
          ),
        );
      } else {
        chatPath = this.path(
          isDalle3 ? OpenaiPath.ImagePath : OpenaiPath.ChatPath,
        );
      }
      if (shouldStream) {
        let index = -1;
        const [tools, funcs] = usePluginStore
          .getState()
          .getAsTools(
            useChatStore.getState().currentSession().mask?.plugin || [],
          );
        // console.log("getAsTools", tools, funcs);
        stream(
          chatPath,
          requestPayload,
          getHeaders(),
          tools as any,
          funcs,
          controller,
          // parseSSE
          (text: string, runTools: ChatMessageTool[]) => {
            // console.log("parseSSE", text, runTools);
            const json = JSON.parse(text);
            const choices = json.choices as Array<{
              delta: {
                content: string;
                tool_calls: ChatMessageTool[];
              };
            }>;
            const tool_calls = choices[0]?.delta?.tool_calls;
            if (tool_calls?.length > 0) {
              const id = tool_calls[0]?.id;
              const args = tool_calls[0]?.function?.arguments;
              if (id) {
                index += 1;
                runTools.push({
                  id,
                  type: tool_calls[0]?.type,
                  function: {
                    name: tool_calls[0]?.function?.name as string,
                    arguments: args,
                  },
                });
              } else {
                // @ts-ignore
                runTools[index]["function"]["arguments"] += args;
              }
            }
            return choices[0]?.delta?.content;
          },
          // processToolMessage, include tool_calls message and tool call results
          (
            requestPayload: RequestPayload,
            toolCallMessage: any,
            toolCallResult: any[],
          ) => {
            // reset index value
            index = -1;
            // @ts-ignore
            requestPayload?.messages?.splice(
              // @ts-ignore
              requestPayload?.messages?.length,
              0,
              toolCallMessage,
              ...toolCallResult,
            );
          },
          options,
        );
      } else {
        const chatPayload = {
          method: "POST",
          body: JSON.stringify(requestPayload),
          signal: controller.signal,
          headers: getHeaders(),
        };

        // make a fetch request
        const requestTimeoutId = setTimeout(
          () => controller.abort(),
          isDalle3 || isO1 ? REQUEST_TIMEOUT_MS * 4 : REQUEST_TIMEOUT_MS, // dalle3 using b64_json is slow.
        );

        const res = await fetch(chatPath, chatPayload);
        clearTimeout(requestTimeoutId);

        const resJson = await res.json();
        const message = await this.extractMessage(resJson);
        options.onFinish(message, res);
      }
    } catch (e) {
      console.log("[Request] failed to make a chat request", e);
      options.onError?.(e as Error);
    }
  }
  async usage() {
    const formatDate = (d: Date) =>
      `${d.getFullYear()}-${(d.getMonth() + 1).toString().padStart(2, "0")}-${d
        .getDate()
        .toString()
        .padStart(2, "0")}`;
    const ONE_DAY = 1 * 24 * 60 * 60 * 1000;
    const now = new Date();
    const startOfMonth = new Date(now.getFullYear(), now.getMonth(), 1);
    const startDate = formatDate(startOfMonth);
    const endDate = formatDate(new Date(Date.now() + ONE_DAY));

    const [used, subs] = await Promise.all([
      fetch(
        this.path(
          `${OpenaiPath.UsagePath}?start_date=${startDate}&end_date=${endDate}`,
        ),
        {
          method: "GET",
          headers: getHeaders(),
        },
      ),
      fetch(this.path(OpenaiPath.SubsPath), {
        method: "GET",
        headers: getHeaders(),
      }),
    ]);

    if (used.status === 401) {
      throw new Error(Locale.Error.Unauthorized);
    }

    if (!used.ok || !subs.ok) {
      throw new Error("Failed to query usage from openai");
    }

    const response = (await used.json()) as {
      total_usage?: number;
      error?: {
        type: string;
        message: string;
      };
    };

    const total = (await subs.json()) as {
      hard_limit_usd?: number;
    };

    if (response.error && response.error.type) {
      throw Error(response.error.message);
    }

    if (response.total_usage) {
      response.total_usage = Math.round(response.total_usage) / 100;
    }

    if (total.hard_limit_usd) {
      total.hard_limit_usd = Math.round(total.hard_limit_usd * 100) / 100;
    }

    return {
      used: response.total_usage,
      total: total.hard_limit_usd,
    } as LLMUsage;
  }

  async models(): Promise<LLMModel[]> {
    if (this.disableListModels) {
      return DEFAULT_MODELS.slice();
    }

    const res = await fetch(this.path(OpenaiPath.ListModelPath), {
      method: "GET",
      headers: {
        ...getHeaders(),
      },
    });

    const resJson = (await res.json()) as OpenAIListModelResponse;
    const chatModels = resJson.data?.filter(
      (m) => m.id.startsWith("gpt-") || m.id.startsWith("chatgpt-"),
    );
    console.log("[Models]", chatModels);

    if (!chatModels) {
      return [];
    }

    //由于目前 OpenAI 的 disableListModels 默认为 true，所以当前实际不会运行到这场
    let seq = 1000; //同 Constant.ts 中的排序保持一致
    return chatModels.map((m) => ({
      name: m.id,
      available: true,
      sorted: seq++,
      provider: {
        id: "openai",
        providerName: "OpenAI",
        providerType: "openai",
        sorted: 1,
      },
    }));
  }
}
export { OpenaiPath };
-												fix: fix gemini issue when using app (#4013)

* chore: update path

* fix: fix google auth logic

* fix: not using header authorization for google api

* chore: revert to allow stream
											
										
										
											2024-02-07 14:17:11 +09:00
+								"use client";
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								// azure and openai, using same models. so using same LLMApi.
-												feat: close #2175 use default api host if endpoint is empty

											
										
										
											2023-06-29 00:12:35 +09:00
+								import {
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
+								  ApiPath,
-												remove DEFAULT_API_HOST

											
										
										
											2024-09-30 02:19:20 +09:00
+								  OPENAI_BASE_URL,
-												feat: #2330 disable /list/models

											
										
										
											2023-07-11 00:19:43 +09:00
+								  DEFAULT_MODELS,
-												feat: close #2175 use default api host if endpoint is empty

											
										
										
											2023-06-29 00:12:35 +09:00
+								  OpenaiPath,
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								  Azure,
-												feat: close #2175 use default api host if endpoint is empty

											
										
										
											2023-06-29 00:12:35 +09:00
+								  REQUEST_TIMEOUT_MS,
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
+								  ServiceProvider,
-												feat: close #2175 use default api host if endpoint is empty

											
										
										
											2023-06-29 00:12:35 +09:00
+								} from "@/app/constant";
-												ts error

											
										
										
											2024-08-29 01:21:26 +09:00
+								import {
 								  ChatMessageTool,
 								  useAccessStore,
 								  useAppConfig,
 								  useChatStore,
-												stash code

											
										
										
											2024-08-29 20:55:09 +09:00
+								  usePluginStore,
-												ts error

											
										
										
											2024-08-29 01:21:26 +09:00
+								} from "@/app/store";
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								import { collectModelsWithDefaultModel } from "@/app/utils/model";
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								import {
 								  preProcessImageContent,
 								  uploadImage,
 								  base64Image2Blob,
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								  stream,
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								} from "@/app/utils/chat";
-												support cloudflare ai gateway

											
										
										
											2024-07-12 13:00:25 +09:00
+								import { cloudflareAIGatewayUrl } from "@/app/utils/cloudflare";
-												dall-e-3 adds 'quality' and 'style' options

											
										
										
											2024-08-10 12:09:07 +09:00
+								import { DalleSize, DalleQuality, DalleStyle } from "@/app/typing";
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
-												Add vision support (#4076)


											
										
										
											2024-02-20 19:04:32 +09:00
+								import {
 								  ChatOptions,
 								  getHeaders,
 								  LLMApi,
 								  LLMModel,
 								  LLMUsage,
 								  MultimodalContent,
-												feat: add tts stt

											
										
										
											2024-08-27 17:21:02 +09:00
+								  SpeechOptions,
-												Add vision support (#4076)


											
										
										
											2024-02-20 19:04:32 +09:00
+								} from "../api";
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								import Locale from "../../locales";
-												feat: close #2621 use better default api url

											
										
										
											2023-08-14 22:36:29 +09:00
+								import { getClientConfig } from "@/app/config/client";
-												Add vision support (#4076)


											
										
										
											2024-02-20 19:04:32 +09:00
+								import {
 								  getMessageTextContent,
 								  isVisionModel,
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								  isDalle3 as _isDalle3,
-												Add vision support (#4076)


											
										
										
											2024-02-20 19:04:32 +09:00
+								} from "@/app/utils";
-												fix: use tauri fetch

											
										
										
											2024-10-16 22:57:07 +09:00
+								import { fetch } from "@/app/utils/stream";
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
-												feat: close #2192 use /list/models to get model ids

											
										
										
											2023-07-05 00:16:24 +09:00
+								export interface OpenAIListModelResponse {
 								  object: string;
 								  data: Array<{
 								    id: string;
 								    object: string;
 								    root: string;
 								  }>;
 								}
-												feat: qwen

											
										
										
											2024-07-07 22:59:56 +09:00
+								export interface RequestPayload {
-												feat: fix no max_tokens in payload when calling openai vision model

											
										
										
											2024-04-08 19:29:08 +09:00
+								  messages: {
 								    role: "system" | "user" | "assistant";
 								    content: string | MultimodalContent[];
 								  }[];
 								  stream?: boolean;
 								  model: string;
 								  temperature: number;
 								  presence_penalty: number;
 								  frequency_penalty: number;
 								  top_p: number;
 								  max_tokens?: number;
-												chore: o1模型使用max_completion_tokens

											
										
										
											2024-11-07 20:45:27 +09:00
+								  max_completion_tokens?: number;
-												feat: fix no max_tokens in payload when calling openai vision model

											
										
										
											2024-04-08 19:29:08 +09:00
+								}
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								export interface DalleRequestPayload {
 								  model: string;
 								  prompt: string;
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								  response_format: "url" | "b64_json";
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								  n: number;
-												fix typescript error

											
										
										
											2024-08-02 19:50:48 +09:00
+								  size: DalleSize;
-												dall-e-3 adds 'quality' and 'style' options

											
										
										
											2024-08-10 12:09:07 +09:00
+								  quality: DalleQuality;
 								  style: DalleStyle;
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								}
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								export class ChatGPTApi implements LLMApi {
-												feat: #2330 disable /list/models

											
										
										
											2023-07-11 00:19:43 +09:00
+								  private disableListModels = true;
-												fix: error

											
										
										
											2024-09-18 16:37:21 +09:00
+								  path(path: string): string {
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
+								    const accessStore = useAccessStore.getState();
-												feat: close #2621 use better default api url

											
										
										
											2023-08-14 22:36:29 +09:00
-												feat: (1) fix issues/4335 and issues/4518

											
										
										
											2024-04-16 15:50:48 +09:00
+								    let baseUrl = "";
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
-												remove makeAzurePath

											
										
										
											2024-07-05 21:15:56 +09:00
+								    const isAzure = path.includes("deployments");
-												feat: (1) fix issues/4335 and issues/4518

											
										
										
											2024-04-16 15:50:48 +09:00
+								    if (accessStore.useCustomConfig) {
 								      if (isAzure && !accessStore.isValidAzure()) {
 								        throw Error(
 								          "incomplete azure config, please check it in your settings page",
 								        );
 								      }
 								      baseUrl = isAzure ? accessStore.azureUrl : accessStore.openaiUrl;
 								    }
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
 								    if (baseUrl.length === 0) {
-												feat: close #2621 use better default api url

											
										
										
											2023-08-14 22:36:29 +09:00
+								      const isApp = !!getClientConfig()?.isApp;
-												remove makeAzurePath

											
										
										
											2024-07-05 21:15:56 +09:00
+								      const apiPath = isAzure ? ApiPath.Azure : ApiPath.OpenAI;
-												update

											
										
										
											2024-09-30 02:44:27 +09:00
+								      baseUrl = isApp ? OPENAI_BASE_URL : apiPath;
-												feat: close #2175 use default api host if endpoint is empty

											
										
										
											2023-06-29 00:12:35 +09:00
+								    }
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
 								    if (baseUrl.endsWith("/")) {
 								      baseUrl = baseUrl.slice(0, baseUrl.length - 1);
 								    }
-												remove makeAzurePath

											
										
										
											2024-07-05 21:15:56 +09:00
+								    if (
 								      !baseUrl.startsWith("http") &&
 								      !isAzure &&
 								      !baseUrl.startsWith(ApiPath.OpenAI)
 								    ) {
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
+								      baseUrl = "https://" + baseUrl;
-												fix: #1509 openai url split

											
										
										
											2023-05-16 01:22:11 +09:00
+								    }
-												feat: close #935 add azure support

											
										
										
											2023-11-10 03:43:30 +09:00
-												fix: fix gemini issue when using app (#4013)

* chore: update path

* fix: fix google auth logic

* fix: not using header authorization for google api

* chore: revert to allow stream
											
										
										
											2024-02-07 14:17:11 +09:00
+								    console.log("[Proxy Endpoint] ", baseUrl, path);
-												support cloudflare ai gateway

											
										
										
											2024-07-12 13:00:25 +09:00
+								    // try rebuild url, when using cloudflare ai gateway in client
 								    return cloudflareAIGatewayUrl([baseUrl, path].join("/"));
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								  }
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								  async extractMessage(res: any) {
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								    if (res.error) {
 								      return "```\n" + JSON.stringify(res, null, 4) + "\n```";
 								    }
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								    // dalle3 model return url, using url create image message
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								    if (res.data) {
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								      let url = res.data?.at(0)?.url ?? "";
 								      const b64_json = res.data?.at(0)?.b64_json ?? "";
 								      if (!url && b64_json) {
 								        // uploadImage
 								        url = await uploadImage(base64Image2Blob(b64_json, "image/png"));
 								      }
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								      return [
 								        {
 								          type: "image_url",
 								          image_url: {
 								            url,
 								          },
 								        },
 								      ];
 								    }
-												fix: empty response

											
										
										
											2024-08-02 23:16:08 +09:00
+								    return res.choices?.at(0)?.message?.content ?? res;
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								  }
-												feat: add tts stt

											
										
										
											2024-08-27 17:21:02 +09:00
+								  async speech(options: SpeechOptions): Promise<ArrayBuffer> {
 								    const requestPayload = {
 								      model: options.model,
 								      input: options.input,
 								      voice: options.voice,
 								      response_format: options.response_format,
 								      speed: options.speed,
 								    };
 								    console.log("[Request] openai speech payload: ", requestPayload);
 								    const controller = new AbortController();
 								    options.onController?.(controller);
 								    try {
-												fix: error

											
										
										
											2024-09-18 16:37:21 +09:00
+								      const speechPath = this.path(OpenaiPath.SpeechPath);
-												feat: add tts stt

											
										
										
											2024-08-27 17:21:02 +09:00
+								      const speechPayload = {
 								        method: "POST",
 								        body: JSON.stringify(requestPayload),
 								        signal: controller.signal,
 								        headers: getHeaders(),
 								      };
 								      // make a fetch request
 								      const requestTimeoutId = setTimeout(
 								        () => controller.abort(),
 								        REQUEST_TIMEOUT_MS,
 								      );
 								      const res = await fetch(speechPath, speechPayload);
 								      clearTimeout(requestTimeoutId);
 								      return await res.arrayBuffer();
 								    } catch (e) {
 								      console.log("[Request] failed to make a speech request", e);
 								      throw e;
 								    }
 								  }
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								  async chat(options: ChatOptions) {
 								    const modelConfig = {
 								      ...useAppConfig.getState().modelConfig,
 								      ...useChatStore.getState().currentSession().mask.modelConfig,
 								      ...{
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								        model: options.config.model,
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								        providerName: options.config.providerName,
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								      },
 								    };
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								    let requestPayload: RequestPayload | DalleRequestPayload;
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								    const isDalle3 = _isDalle3(options.config.model);
-												feat: add o1 model

											
										
										
											2024-09-13 14:18:07 +09:00
+								    const isO1 = options.config.model.startsWith("o1");
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								    if (isDalle3) {
-												fix typescript error

											
										
										
											2024-08-02 19:50:48 +09:00
+								      const prompt = getMessageTextContent(
 								        options.messages.slice(-1)?.pop() as any,
 								      );
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								      requestPayload = {
 								        model: options.config.model,
 								        prompt,
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								        // URLs are only valid for 60 minutes after the image has been generated.
 								        response_format: "b64_json", // using b64_json, and save image in CacheStorage
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								        n: 1,
 								        size: options.config?.size ?? "1024x1024",
-												dall-e-3 adds 'quality' and 'style' options

											
										
										
											2024-08-10 12:09:07 +09:00
+								        quality: options.config?.quality ?? "standard",
 								        style: options.config?.style ?? "vivid",
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								      };
 								    } else {
 								      const visionModel = isVisionModel(options.config.model);
 								      const messages: ChatOptions["messages"] = [];
 								      for (const v of options.messages) {
 								        const content = visionModel
 								          ? await preProcessImageContent(v.content)
 								          : getMessageTextContent(v);
-												fix: give o1 some time to think twice

											
										
										
											2024-09-13 17:25:04 +09:00
+								        if (!(isO1 && v.role === "system"))
-												feat: add o1 model

											
										
										
											2024-09-13 14:18:07 +09:00
+								          messages.push({ role: v.role, content });
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								      }
-												feat: add o1 model

											
										
										
											2024-09-13 14:18:07 +09:00
+								      // O1 not support image, tools (plugin in ChatGPTNextWeb) and system, stream, logprobs, temperature, top_p, n, presence_penalty, frequency_penalty yet.
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								      requestPayload = {
 								        messages,
-												use stream when request o1

											
										
										
											2024-11-21 12:46:10 +09:00
+								        stream: options.config.stream,
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								        model: modelConfig.model,
-												feat: add o1 model

											
										
										
											2024-09-13 14:18:07 +09:00
+								        temperature: !isO1 ? modelConfig.temperature : 1,
 								        presence_penalty: !isO1 ? modelConfig.presence_penalty : 0,
 								        frequency_penalty: !isO1 ? modelConfig.frequency_penalty : 0,
 								        top_p: !isO1 ? modelConfig.top_p : 1,
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								        // max_tokens: Math.max(modelConfig.max_tokens, 1024),
 								        // Please do not ask me why not send max_tokens, no reason, this param is just shit, I dont want to explain anymore.
 								      };
-												chore: o1模型使用max_completion_tokens

											
										
										
											2024-11-07 20:45:27 +09:00
+								      // O1 使用 max_completion_tokens 控制token数 (https://platform.openai.com/docs/guides/reasoning#controlling-costs)
 								      if (isO1) {
 								        requestPayload["max_completion_tokens"] = modelConfig.max_tokens;
 								      }
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								      // add max_tokens to vision model
-												fix: remove the visual model judgment method that checks if the model name contains 'preview' from the openai api to prevent models like o1-preview from being classified as visual models

											
										
										
											2024-09-13 13:56:28 +09:00
+								      if (visionModel) {
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								        requestPayload["max_tokens"] = Math.max(modelConfig.max_tokens, 4000);
 								      }
-												fix: add max_tokens when using vision model (#4157)


											
										
										
											2024-02-27 18:28:01 +09:00
+								    }
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								    console.log("[Request] openai payload: ", requestPayload);
-												use stream when request o1

											
										
										
											2024-11-21 12:46:10 +09:00
+								    const shouldStream = !isDalle3 && !!options.config.stream;
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								    const controller = new AbortController();
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								    options.onController?.(controller);
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
 								    try {
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								      let chatPath = "";
-												chore: optimize the code

											
										
										
											2024-07-06 00:56:10 +09:00
+								      if (modelConfig.providerName === ServiceProvider.Azure) {
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								        // find model, and get displayName as deployName
 								        const { models: configModels, customModels: configCustomModels } =
 								          useAppConfig.getState();
-												using default azure api-version value

											
										
										
											2024-07-06 01:05:59 +09:00
+								        const {
 								          defaultModel,
 								          customModels: accessCustomModels,
 								          useCustomConfig,
 								        } = useAccessStore.getState();
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								        const models = collectModelsWithDefaultModel(
 								          configModels,
 								          [configCustomModels, accessCustomModels].join(","),
 								          defaultModel,
 								        );
 								        const model = models.find(
 								          (model) =>
-												chore: optimize the code

											
										
										
											2024-07-06 00:56:10 +09:00
+								            model.name === modelConfig.model &&
 								            model?.provider?.providerName === ServiceProvider.Azure,
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								        );
-												remove makeAzurePath

											
										
										
											2024-07-05 21:15:56 +09:00
+								        chatPath = this.path(
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								          (isDalle3 ? Azure.ImagePath : Azure.ChatPath)(
-												fix ts

											
										
										
											2024-07-05 21:20:21 +09:00
+								            (model?.displayName ?? model?.name) as string,
-												using default azure api-version value

											
										
										
											2024-07-06 01:05:59 +09:00
+								            useCustomConfig ? useAccessStore.getState().azureApiVersion : "",
-												remove makeAzurePath

											
										
										
											2024-07-05 21:15:56 +09:00
+								          ),
 								        );
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								      } else {
-												add dalle3 model

											
										
										
											2024-08-02 19:00:42 +09:00
+								        chatPath = this.path(
 								          isDalle3 ? OpenaiPath.ImagePath : OpenaiPath.ChatPath,
 								        );
-												support azure deployment name

											
										
										
											2024-07-05 20:59:45 +09:00
+								      }
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								      if (shouldStream) {
-												hotfix openai function call tool_calls no index

											
										
										
											2024-09-22 19:53:51 +09:00
+								        let index = -1;
-												stash code

											
										
										
											2024-08-30 18:31:20 +09:00
+								        const [tools, funcs] = usePluginStore
-												stash code

											
										
										
											2024-08-29 20:55:09 +09:00
+								          .getState()
-												ts error

											
										
										
											2024-08-31 00:39:08 +09:00
+								          .getAsTools(
-												fix(#5378): default plugin ids to empty array

											
										
										
											2024-09-07 22:32:18 +09:00
+								            useChatStore.getState().currentSession().mask?.plugin || [],
-												ts error

											
										
										
											2024-08-31 00:39:08 +09:00
+								          );
-												claude support function call

											
										
										
											2024-09-02 22:45:47 +09:00
+								        // console.log("getAsTools", tools, funcs);
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								        stream(
 								          chatPath,
 								          requestPayload,
 								          getHeaders(),
-												ts error

											
										
										
											2024-08-31 00:39:08 +09:00
+								          tools as any,
-												stash code

											
										
										
											2024-08-30 18:31:20 +09:00
+								          funcs,
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								          controller,
-												add processToolMessage callback

											
										
										
											2024-08-29 18:28:15 +09:00
+								          // parseSSE
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								          (text: string, runTools: ChatMessageTool[]) => {
-												add processToolMessage callback

											
										
										
											2024-08-29 18:28:15 +09:00
+								            // console.log("parseSSE", text, runTools);
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								            const json = JSON.parse(text);
 								            const choices = json.choices as Array<{
 								              delta: {
 								                content: string;
 								                tool_calls: ChatMessageTool[];
-												stash code

											
										
										
											2024-08-29 00:58:46 +09:00
+								              };
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								            }>;
 								            const tool_calls = choices[0]?.delta?.tool_calls;
 								            if (tool_calls?.length > 0) {
 								              const id = tool_calls[0]?.id;
 								              const args = tool_calls[0]?.function?.arguments;
 								              if (id) {
-												hotfix openai function call tool_calls no index

											
										
										
											2024-09-22 19:59:49 +09:00
+								                index += 1;
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								                runTools.push({
 								                  id,
 								                  type: tool_calls[0]?.type,
-												stash code

											
										
										
											2024-08-29 00:58:46 +09:00
+								                  function: {
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								                    name: tool_calls[0]?.function?.name as string,
 								                    arguments: args,
-												stash code

											
										
										
											2024-08-29 00:58:46 +09:00
+								                  },
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								                });
 								              } else {
 								                // @ts-ignore
 								                runTools[index]["function"]["arguments"] += args;
-												fixup: add more error info

											
										
										
											2023-05-16 02:58:58 +09:00
+								              }
-												fix: #1498 missing text caused by streaming

											
										
										
											2023-05-16 02:25:16 +09:00
+								            }
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								            return choices[0]?.delta?.content;
-												fix: #1498 missing text caused by streaming

											
										
										
											2023-05-16 02:25:16 +09:00
+								          },
-												add processToolMessage callback

											
										
										
											2024-08-29 18:28:15 +09:00
+								          // processToolMessage, include tool_calls message and tool call results
 								          (
 								            requestPayload: RequestPayload,
 								            toolCallMessage: any,
 								            toolCallResult: any[],
 								          ) => {
-												hotfix openai function call tool_calls no index

											
										
										
											2024-09-22 19:53:51 +09:00
+								            // reset index value
 								            index = -1;
-												add processToolMessage callback

											
										
										
											2024-08-29 18:28:15 +09:00
+								            // @ts-ignore
 								            requestPayload?.messages?.splice(
 								              // @ts-ignore
 								              requestPayload?.messages?.length,
 ,
 								              toolCallMessage,
 								              ...toolCallResult,
 								            );
-												fix: #1498 missing text caused by streaming

											
										
										
											2023-05-16 02:25:16 +09:00
+								          },
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								          options,
 								        );
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								      } else {
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								        const chatPayload = {
 								          method: "POST",
 								          body: JSON.stringify(requestPayload),
 								          signal: controller.signal,
 								          headers: getHeaders(),
 								        };
-												stash code

											
										
										
											2024-08-29 00:58:46 +09:00
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								        // make a fetch request
 								        const requestTimeoutId = setTimeout(
 								          () => controller.abort(),
-												仅修改o1的超时时间为4分钟，减少o1系列模型请求失败的情况

											
										
										
											2024-10-14 17:31:17 +09:00
+								          isDalle3 || isO1 ? REQUEST_TIMEOUT_MS * 4 : REQUEST_TIMEOUT_MS, // dalle3 using b64_json is slow.
-												create common function stream for fetchEventSource

											
										
										
											2024-08-29 18:14:23 +09:00
+								        );
-												stash code

											
										
										
											2024-08-29 00:58:46 +09:00
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								        const res = await fetch(chatPath, chatPayload);
-												fix: typo reqestTimeoutId -> requestTimeoutId
											
										
										
											2023-05-16 09:59:30 +09:00
+								        clearTimeout(requestTimeoutId);
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
 								        const resJson = await res.json();
-												using b64_json for dall-e-3

											
										
										
											2024-08-02 21:58:21 +09:00
+								        const message = await this.extractMessage(resJson);
-												fix: onfinish responseRes

											
										
										
											2024-11-04 18:00:45 +09:00
+								        options.onFinish(message, res);
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								      }
 								    } catch (e) {
-												typo fix
											
										
										
											2023-08-01 11:16:36 +09:00
+								      console.log("[Request] failed to make a chat request", e);
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								      options.onError?.(e as Error);
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								    }
 								  }
 								  async usage() {
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								    const formatDate = (d: Date) =>
 								      `${d.getFullYear()}-${(d.getMonth() + 1).toString().padStart(2, "0")}-${d
 								        .getDate()
 								        .toString()
 								        .padStart(2, "0")}`;
 								    const ONE_DAY = 1 * 24 * 60 * 60 * 1000;
 								    const now = new Date();
 								    const startOfMonth = new Date(now.getFullYear(), now.getMonth(), 1);
 								    const startDate = formatDate(startOfMonth);
 								    const endDate = formatDate(new Date(Date.now() + ONE_DAY));
 								    const [used, subs] = await Promise.all([
 								      fetch(
 								        this.path(
-												feat: white url list for openai security

											
										
										
											2023-06-13 01:39:29 +09:00
+								          `${OpenaiPath.UsagePath}?start_date=${startDate}&end_date=${endDate}`,
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								        ),
 								        {
 								          method: "GET",
 								          headers: getHeaders(),
 								        },
 								      ),
-												feat: white url list for openai security

											
										
										
											2023-06-13 01:39:29 +09:00
+								      fetch(this.path(OpenaiPath.SubsPath), {
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								        method: "GET",
 								        headers: getHeaders(),
 								      }),
 								    ]);
-												fix: #1611 show corret message when can not query usage

											
										
										
											2023-05-19 01:27:25 +09:00
+								    if (used.status === 401) {
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								      throw new Error(Locale.Error.Unauthorized);
 								    }
-												fix: #1611 show corret message when can not query usage

											
										
										
											2023-05-19 01:27:25 +09:00
+								    if (!used.ok || !subs.ok) {
 								      throw new Error("Failed to query usage from openai");
 								    }
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								    const response = (await used.json()) as {
 								      total_usage?: number;
 								      error?: {
 								        type: string;
 								        message: string;
 								      };
 								    };
 								    const total = (await subs.json()) as {
 								      hard_limit_usd?: number;
 								    };
 								    if (response.error && response.error.type) {
 								      throw Error(response.error.message);
 								    }
 								    if (response.total_usage) {
 								      response.total_usage = Math.round(response.total_usage) / 100;
 								    }
 								    if (total.hard_limit_usd) {
 								      total.hard_limit_usd = Math.round(total.hard_limit_usd * 100) / 100;
 								    }
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								    return {
-												refactor: #1000 #1179 api layer for client-side only mode and local models

											
										
										
											2023-05-15 02:33:46 +09:00
+								      used: response.total_usage,
 								      total: total.hard_limit_usd,
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								    } as LLMUsage;
 								  }
-												feat: close #2192 use /list/models to get model ids

											
										
										
											2023-07-05 00:16:24 +09:00
 								  async models(): Promise<LLMModel[]> {
-												feat: #2330 disable /list/models

											
										
										
											2023-07-11 00:19:43 +09:00
+								    if (this.disableListModels) {
 								      return DEFAULT_MODELS.slice();
 								    }
-												feat: close #2192 use /list/models to get model ids

											
										
										
											2023-07-05 00:16:24 +09:00
+								    const res = await fetch(this.path(OpenaiPath.ListModelPath), {
 								      method: "GET",
 								      headers: {
 								        ...getHeaders(),
 								      },
 								    });
 								    const resJson = (await res.json()) as OpenAIListModelResponse;
-												add chatgpt-4o-latest

											
										
										
											2024-09-07 02:42:56 +09:00
+								    const chatModels = resJson.data?.filter(
 								      (m) => m.id.startsWith("gpt-") || m.id.startsWith("chatgpt-"),
 								    );
-												feat: close #2192 use /list/models to get model ids

											
										
										
											2023-07-05 00:16:24 +09:00
+								    console.log("[Models]", chatModels);
-												fix: #2280 auto-detect models from 'list/models'

											
										
										
											2023-07-09 19:03:06 +09:00
+								    if (!chatModels) {
 								      return [];
 								    }
-												🐛 fix(openai): 上次 commit 后 openai.ts 文件中出现类型不匹配的 bug

											
										
										
											2024-08-05 21:26:48 +09:00
+								    //由于目前 OpenAI 的 disableListModels 默认为 true，所以当前实际不会运行到这场
 								    let seq = 1000; //同 Constant.ts 中的排序保持一致
-												fix: #2280 auto-detect models from 'list/models'

											
										
										
											2023-07-09 19:03:06 +09:00
+								    return chatModels.map((m) => ({
 								      name: m.id,
 								      available: true,
-												🐛 fix(openai): 上次 commit 后 openai.ts 文件中出现类型不匹配的 bug

											
										
										
											2024-08-05 21:26:48 +09:00
+								      sorted: seq++,
-												fix: fix type errors

											
										
										
											2023-12-24 03:39:06 +09:00
+								      provider: {
 								        id: "openai",
 								        providerName: "OpenAI",
 								        providerType: "openai",
-												🐛 fix(openai): 上次 commit 后 openai.ts 文件中出现类型不匹配的 bug

											
										
										
											2024-08-05 21:26:48 +09:00
+								        sorted: 1,
-												fix: fix type errors

											
										
										
											2023-12-24 03:39:06 +09:00
+								      },
-												fix: #2280 auto-detect models from 'list/models'

											
										
										
											2023-07-09 19:03:06 +09:00
+								    }));
-												feat: close #2192 use /list/models to get model ids

											
										
										
											2023-07-05 00:16:24 +09:00
+								  }
-												refactor: llm client api

											
										
										
											2023-05-15 00:00:17 +09:00
+								}
-												feat: white url list for openai security

											
										
										
											2023-06-13 01:39:29 +09:00
+								export { OpenaiPath };