Reorganized the solution structure to align with a layered architecture: - Replaced `src` folder with `core`, `infrastructure`, and `presentation`. - Moved projects to their respective folders. - Added `DigitalData.MessagingService.Publisher.Abstraction` project. - Removed `DigitalData.MessagingService.Client` project. Updated project configurations and nesting in the solution file. Added `appsettings.Secrets.json` with RabbitMQ and email account settings: - RabbitMQ configuration includes hostname, port, credentials, and queue/exchange details. - Email configuration includes SMTP server details and credentials.
96 lines
3.2 KiB
C#
96 lines
3.2 KiB
C#
using DevExpress.Pdf;
|
|
using DigitalData.MessagingService.Application.Common.Interfaces;
|
|
using DigitalData.MessagingService.Domain.Exceptions;
|
|
|
|
namespace DigitalData.MessagingService.Infrastructure.Services;
|
|
|
|
/// <summary>
|
|
/// PDF processing service using DevExpress.Pdf.
|
|
/// Implements PDF validation and embedded file extraction using streams.
|
|
/// </summary>
|
|
public class DevExpressPdfProcessingService : IPdfProcessingService
|
|
{
|
|
public Task<bool> ValidatePdfAsync(Stream pdfStream, CancellationToken cancellationToken = default)
|
|
{
|
|
ArgumentNullException.ThrowIfNull(pdfStream);
|
|
|
|
if (!pdfStream.CanRead)
|
|
throw new ArgumentException("Stream must be readable.", nameof(pdfStream));
|
|
|
|
if (!pdfStream.CanSeek)
|
|
throw new ArgumentException("Stream must be seekable.", nameof(pdfStream));
|
|
|
|
if (pdfStream.Position != 0)
|
|
pdfStream.Position = 0;
|
|
|
|
using var processor = new PdfDocumentProcessor();
|
|
processor.LoadDocument(pdfStream);
|
|
|
|
return Task.FromResult(true);
|
|
}
|
|
|
|
public async Task<IEnumerable<string>> ExtractEmbeddedFilesAsync(
|
|
Stream pdfStream,
|
|
string outputDirectory,
|
|
CancellationToken cancellationToken = default)
|
|
{
|
|
ArgumentNullException.ThrowIfNull(pdfStream);
|
|
ArgumentException.ThrowIfNullOrWhiteSpace(outputDirectory);
|
|
|
|
if (!pdfStream.CanRead)
|
|
throw new ArgumentException("Stream must be readable.", nameof(pdfStream));
|
|
|
|
if (!pdfStream.CanSeek)
|
|
throw new ArgumentException("Stream must be seekable.", nameof(pdfStream));
|
|
|
|
if (pdfStream.Position != 0)
|
|
pdfStream.Position = 0;
|
|
|
|
if (!Directory.Exists(outputDirectory))
|
|
Directory.CreateDirectory(outputDirectory);
|
|
|
|
using var processor = new PdfDocumentProcessor();
|
|
processor.LoadDocument(pdfStream);
|
|
|
|
var extractedFiles = new List<string>();
|
|
var attachments = processor.Document.FileAttachments;
|
|
|
|
if (attachments == null || !attachments.Any())
|
|
return extractedFiles;
|
|
|
|
foreach (var attachment in attachments)
|
|
{
|
|
var fileName = attachment.FileName ?? $"attachment_{Guid.NewGuid()}.dat";
|
|
var outputPath = Path.Combine(outputDirectory, fileName);
|
|
|
|
var fileData = attachment.Data;
|
|
if (fileData == null || fileData.Length == 0)
|
|
continue;
|
|
|
|
await File.WriteAllBytesAsync(outputPath, fileData, cancellationToken);
|
|
extractedFiles.Add(outputPath);
|
|
}
|
|
|
|
return extractedFiles;
|
|
}
|
|
|
|
public Task<int> GetPageCountAsync(Stream pdfStream, CancellationToken cancellationToken = default)
|
|
{
|
|
ArgumentNullException.ThrowIfNull(pdfStream);
|
|
|
|
if (!pdfStream.CanRead)
|
|
throw new ArgumentException("Stream must be readable.", nameof(pdfStream));
|
|
|
|
if (!pdfStream.CanSeek)
|
|
throw new ArgumentException("Stream must be seekable.", nameof(pdfStream));
|
|
|
|
if (pdfStream.Position != 0)
|
|
pdfStream.Position = 0;
|
|
|
|
using var processor = new PdfDocumentProcessor();
|
|
processor.LoadDocument(pdfStream);
|
|
|
|
return Task.FromResult(processor.Document.Pages.Count);
|
|
}
|
|
}
|