mirror of
https://github.com/danny-avila/LibreChat.git
synced 2025-12-28 14:18:51 +01:00
Some checks failed
Docker Dev Branch Images Build / build (Dockerfile, lc-dev, node) (push) Waiting to run
Docker Dev Branch Images Build / build (Dockerfile.multi, lc-dev-api, api-build) (push) Waiting to run
Docker Dev Images Build / build (Dockerfile, librechat-dev, node) (push) Has been cancelled
Docker Dev Images Build / build (Dockerfile.multi, librechat-dev-api, api-build) (push) Has been cancelled
Sync Locize Translations & Create Translation PR / Sync Translation Keys with Locize (push) Has been cancelled
Sync Locize Translations & Create Translation PR / Create Translation PR on Version Published (push) Has been cancelled
* 📎 feat: Direct Provider Attachment Support for Multimodal Content * 📑 feat: Anthropic Direct Provider Upload (#9072) * feat: implement Anthropic native PDF support with document preservation - Add comprehensive debug logging throughout PDF processing pipeline - Refactor attachment processing to separate image and document handling - Create distinct addImageURLs(), addDocuments(), and processAttachments() methods - Fix critical bugs in stream handling and parameter passing - Add streamToBuffer utility for proper stream-to-buffer conversion - Remove api/agents submodule from repository 🤖 Generated with [Claude Code](https://claude.ai/code) Co-Authored-By: Claude <noreply@anthropic.com> * chore: remove out of scope formatting changes * fix: stop duplication of file in chat on end of response stream * chore: bring back file search and ocr options * chore: localize upload to provider string in file menu * refactor: change createMenuItems args to fit new pattern introduced by anthropic-native-pdf-support * feat: add cache point for pdfs processed by anthropic endpoint since they are unlikely to change and should benefit from caching * feat: combine Upload Image into Upload to Provider since they both perform direct upload and change provider upload icon to reflect multimodal upload * feat: add citations support according to docs * refactor: remove redundant 'document' check since documents are handled properly by formatMessage in the agents repo now * refactor: change upload logic so anthropic endpoint isn't exempted from normal upload path using Agents for consistency with the rest of the upload logic * fix: include width and height in return from uploadLocalFile so images are correctly identified when going through an AgentUpload in addImageURLs * chore: remove client specific handling since the direct provider stuff is handled by the agent client * feat: handle documents in AgentClient so no need for change to agents repo * chore: removed unused changes * chore: remove auto generated comments from OG commit * feat: add logic for agents to use direct to provider uploads if supported (currently just anthropic) * fix: reintroduce role check to fix render error because of undefined value for Content Part * fix: actually fix render bug by using proper isCreatedByUser check and making sure our mutation of formattedMessage.content is consistent --------- Co-authored-by: Andres Restrepo <andres@thelinuxkid.com> Co-authored-by: Claude <noreply@anthropic.com> 📁 feat: Send Attachments Directly to Provider (OpenAI) (#9098) * refactor: change references from direct upload to direct attach to better reflect functionality since we are just using base64 encoding strategy now rather than Files/File API for sending our attachments directly to the provider, the upload nomenclature no longer makes sense. direct_attach better describes the different methods of sending attachments to providers anyways even if we later introduce direct upload support * feat: add upload to provider option for openai (and agent) ui * chore: move anthropic pdf validator over to packages/api * feat: simple pdf validation according to openai docs * feat: add provider agnostic validatePdf logic to start handling multiple endpoints * feat: add handling for openai specific documentPart formatting * refactor: move require statement to proper place at top of file * chore: add in openAI endpoint for the rest of the document handling logic * feat: add direct attach support for azureOpenAI endpoint and agents * feat: add pdf validation for azureOpenAI endpoint * refactor: unify all the endpoint checks with isDocumentSupportedEndpoint * refactor: consolidate Upload to Provider vs Upload image logic for clarity * refactor: remove anthropic from anthropic_multimodal fileType since we support multiple providers now 🗂️ feat: Send Attachments Directly to Provider (Google) (#9100) * feat: add validation for google PDFs and add google endpoint as a document supporting endpoint * feat: add proper pdf formatting for google endpoints (requires PR #14 in agents) * feat: add multimodal support for google endpoint attachments * feat: add audio file svg * fix: refactor attachments logic so multi-attachment messages work properly * feat: add video file svg * fix: allows for followup questions of uploaded multimodal attachments * fix: remove incorrect final message filtering that was breaking Attachment component rendering fix: manualy rename 'documents' to 'Documents' in git since it wasn't picked up due to case insensitivity in dir name fix: add logic so filepicker for a google agent has proper filetype filtering 🛫 refactor: Move Encoding Logic to packages/api (#9182) * refactor: move audio encode over to TS * refactor: audio encoding now functional in LC again * refactor: move video encode over to TS * refactor: move document encode over to TS * refactor: video encoding now functional in LC again * refactor: document encoding now functional in LC again * fix: extend file type options in AttachFileMenu to include 'google_multimodal' and update dependency array to include agent?.provider * feat: only accept pdfs if responses api is enabled for openai convos chore: address ESLint comments chore: add missing audio mimetype * fix: type safety for message content parts and improve null handling * chore: reorder AttachFileMenuProps for consistency and clarity * chore: import order in AttachFileMenu * fix: improve null handling for text parts in parseTextParts function * fix: remove no longer used unsupported capability error message for file uploads * fix: OpenAI Direct File Attachment Format * fix: update encodeAndFormatDocuments to support OpenAI responses API and enhance document result types * refactor: broaden providers supported for documents * feat: enhance DragDrop context and modal to support document uploads based on provider capabilities * fix: reorder import statements for consistency in video encoding module --------- Co-authored-by: Dustin Healy <54083382+dustinhealy@users.noreply.github.com>
108 lines
3.5 KiB
TypeScript
108 lines
3.5 KiB
TypeScript
import { Providers } from '@librechat/agents';
|
|
import { isOpenAILikeProvider, isDocumentSupportedProvider } from 'librechat-data-provider';
|
|
import type { IMongoFile } from '@librechat/data-schemas';
|
|
import type { Request } from 'express';
|
|
import type { StrategyFunctions, DocumentResult } from '~/types/files';
|
|
import { validatePdf } from '~/files/validation';
|
|
import { getFileStream } from './utils';
|
|
|
|
/**
|
|
* Processes and encodes document files for various providers
|
|
* @param req - Express request object
|
|
* @param files - Array of file objects to process
|
|
* @param provider - The provider name
|
|
* @param getStrategyFunctions - Function to get strategy functions
|
|
* @returns Promise that resolves to documents and file metadata
|
|
*/
|
|
export async function encodeAndFormatDocuments(
|
|
req: Request,
|
|
files: IMongoFile[],
|
|
{ provider, useResponsesApi }: { provider: Providers; useResponsesApi?: boolean },
|
|
getStrategyFunctions: (source: string) => StrategyFunctions,
|
|
): Promise<DocumentResult> {
|
|
if (!files?.length) {
|
|
return { documents: [], files: [] };
|
|
}
|
|
|
|
const encodingMethods: Record<string, StrategyFunctions> = {};
|
|
const result: DocumentResult = { documents: [], files: [] };
|
|
|
|
const documentFiles = files.filter(
|
|
(file) => file.type === 'application/pdf' || file.type?.startsWith('application/'),
|
|
);
|
|
|
|
if (!documentFiles.length) {
|
|
return result;
|
|
}
|
|
|
|
const results = await Promise.allSettled(
|
|
documentFiles.map((file) => {
|
|
if (file.type !== 'application/pdf' || !isDocumentSupportedProvider(provider)) {
|
|
return Promise.resolve(null);
|
|
}
|
|
return getFileStream(req, file, encodingMethods, getStrategyFunctions);
|
|
}),
|
|
);
|
|
|
|
for (const settledResult of results) {
|
|
if (settledResult.status === 'rejected') {
|
|
console.error('Document processing failed:', settledResult.reason);
|
|
continue;
|
|
}
|
|
|
|
const processed = settledResult.value;
|
|
if (!processed) continue;
|
|
|
|
const { file, content, metadata } = processed;
|
|
|
|
if (!content || !file) {
|
|
if (metadata) result.files.push(metadata);
|
|
continue;
|
|
}
|
|
|
|
if (file.type === 'application/pdf' && isDocumentSupportedProvider(provider)) {
|
|
const pdfBuffer = Buffer.from(content, 'base64');
|
|
const validation = await validatePdf(pdfBuffer, pdfBuffer.length, provider);
|
|
|
|
if (!validation.isValid) {
|
|
throw new Error(`PDF validation failed: ${validation.error}`);
|
|
}
|
|
|
|
if (provider === Providers.ANTHROPIC) {
|
|
result.documents.push({
|
|
type: 'document',
|
|
source: {
|
|
type: 'base64',
|
|
media_type: 'application/pdf',
|
|
data: content,
|
|
},
|
|
cache_control: { type: 'ephemeral' },
|
|
citations: { enabled: true },
|
|
});
|
|
} else if (useResponsesApi) {
|
|
result.documents.push({
|
|
type: 'input_file',
|
|
filename: file.filename,
|
|
file_data: `data:application/pdf;base64,${content}`,
|
|
});
|
|
} else if (provider === Providers.GOOGLE || provider === Providers.VERTEXAI) {
|
|
result.documents.push({
|
|
type: 'document',
|
|
mimeType: 'application/pdf',
|
|
data: content,
|
|
});
|
|
} else if (isOpenAILikeProvider(provider) && provider != Providers.AZURE) {
|
|
result.documents.push({
|
|
type: 'file',
|
|
file: {
|
|
filename: file.filename,
|
|
file_data: `data:application/pdf;base64,${content}`,
|
|
},
|
|
});
|
|
}
|
|
result.files.push(metadata);
|
|
}
|
|
}
|
|
|
|
return result;
|
|
}
|