2026-02-23 00:51:46 +01:00
|
|
|
|
use bytes::{Bytes, BytesMut};
|
|
|
|
|
|
use futures::{Stream, StreamExt};
|
2026-02-26 00:07:10 +01:00
|
|
|
|
use std::pin::Pin;
|
2026-02-14 01:29:34 +01:00
|
|
|
|
use std::sync::Arc;
|
2025-03-19 19:52:12 +01:00
|
|
|
|
|
|
|
|
|
|
use crate::application::dtos::file_dto::FileDto;
|
2026-05-20 22:56:00 +02:00
|
|
|
|
use crate::application::ports::authorization_ports::AuthorizationEngine;
|
2026-02-08 13:40:23 +01:00
|
|
|
|
use crate::application::ports::file_ports::{FileRetrievalUseCase, OptimizedFileContent};
|
2026-06-26 01:01:44 +02:00
|
|
|
|
use crate::application::ports::resource_access_hook::ResourceAccessHook;
|
2025-03-19 19:52:12 +01:00
|
|
|
|
use crate::application::ports::storage_ports::FileReadPort;
|
|
|
|
|
|
use crate::common::errors::DomainError;
|
2026-05-20 22:56:00 +02:00
|
|
|
|
use crate::domain::services::authorization::{Permission, Resource, Subject};
|
2026-03-03 15:36:42 +00:00
|
|
|
|
use crate::infrastructure::repositories::pg::file_blob_read_repository::FileBlobReadRepository;
|
|
|
|
|
|
use crate::infrastructure::services::file_content_cache::FileContentCache;
|
2026-03-04 23:55:08 +01:00
|
|
|
|
use crate::infrastructure::services::image_transcode_service::{
|
|
|
|
|
|
ImageTranscodeService, OutputFormat,
|
|
|
|
|
|
};
|
2026-05-20 22:56:00 +02:00
|
|
|
|
use crate::infrastructure::services::pg_acl_engine::PgAclEngine;
|
2026-03-04 23:55:08 +01:00
|
|
|
|
use tracing::{debug, info};
|
2026-03-07 14:59:32 +01:00
|
|
|
|
use uuid::Uuid;
|
2026-02-08 13:40:23 +01:00
|
|
|
|
|
|
|
|
|
|
/// Threshold below which files are served from RAM cache (10 MB).
|
|
|
|
|
|
const CACHE_THRESHOLD: u64 = 10 * 1024 * 1024;
|
2025-03-19 19:52:12 +01:00
|
|
|
|
|
2026-02-12 09:41:25 +01:00
|
|
|
|
/// Service for file retrieval operations
|
2026-02-08 13:40:23 +01:00
|
|
|
|
///
|
|
|
|
|
|
/// Implements a multi-tier download strategy:
|
|
|
|
|
|
/// - Tier 0: Write-behind cache (just-uploaded files still in RAM)
|
|
|
|
|
|
/// - Tier 1: Hot cache + optional WebP transcoding (<10 MB)
|
2026-06-22 08:10:54 +00:00
|
|
|
|
/// - Tier 2: Streaming for everything ≥10 MB — CDC chunk reassembly with the
|
|
|
|
|
|
/// backend's read-ahead (`read_prefetch`); no whole-file buffering.
|
2025-03-19 19:52:12 +01:00
|
|
|
|
pub struct FileRetrievalService {
|
2026-03-03 15:36:42 +00:00
|
|
|
|
file_read: Arc<FileBlobReadRepository>,
|
|
|
|
|
|
content_cache: Option<Arc<FileContentCache>>,
|
|
|
|
|
|
transcode: Option<Arc<ImageTranscodeService>>,
|
2026-05-20 22:56:00 +02:00
|
|
|
|
authz: Option<Arc<PgAclEngine>>,
|
2026-06-26 01:01:44 +02:00
|
|
|
|
/// Optional read-event observer. Currently fans out to the Recent-list
|
|
|
|
|
|
/// recorder; future observers (audit trail, "last seen by", …) attach
|
|
|
|
|
|
/// to the same hook so service code only knows the trait, not the impl.
|
|
|
|
|
|
/// `None` for the test/stub path that constructs via [`Self::new`].
|
|
|
|
|
|
resource_access_hook: Option<Arc<dyn ResourceAccessHook>>,
|
2025-03-19 19:52:12 +01:00
|
|
|
|
}
|
|
|
|
|
|
|
|
|
|
|
|
impl FileRetrievalService {
|
2026-05-20 22:56:00 +02:00
|
|
|
|
/// Backward-compatible constructor (simple pass-through). Without the
|
|
|
|
|
|
/// authorization engine, the `*_owned`/`*_with_perms` methods fail closed.
|
|
|
|
|
|
/// Use `new_with_cache` in production.
|
2026-03-03 15:36:42 +00:00
|
|
|
|
pub fn new(file_repository: Arc<FileBlobReadRepository>) -> Self {
|
2026-02-08 13:40:23 +01:00
|
|
|
|
Self {
|
|
|
|
|
|
file_read: file_repository,
|
|
|
|
|
|
content_cache: None,
|
|
|
|
|
|
transcode: None,
|
2026-05-20 22:56:00 +02:00
|
|
|
|
authz: None,
|
2026-06-26 01:01:44 +02:00
|
|
|
|
resource_access_hook: None,
|
2026-02-08 13:40:23 +01:00
|
|
|
|
}
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-20 22:56:00 +02:00
|
|
|
|
/// Constructor for blob-storage model: read + content cache + transcode +
|
|
|
|
|
|
/// ReBAC authorization.
|
2026-02-14 17:54:25 +01:00
|
|
|
|
pub fn new_with_cache(
|
2026-03-03 15:36:42 +00:00
|
|
|
|
file_read: Arc<FileBlobReadRepository>,
|
|
|
|
|
|
content_cache: Arc<FileContentCache>,
|
|
|
|
|
|
transcode: Arc<ImageTranscodeService>,
|
2026-05-20 22:56:00 +02:00
|
|
|
|
authz: Arc<PgAclEngine>,
|
2026-02-14 17:54:25 +01:00
|
|
|
|
) -> Self {
|
|
|
|
|
|
Self {
|
|
|
|
|
|
file_read,
|
|
|
|
|
|
content_cache: Some(content_cache),
|
|
|
|
|
|
transcode: Some(transcode),
|
2026-05-20 22:56:00 +02:00
|
|
|
|
authz: Some(authz),
|
2026-06-26 01:01:44 +02:00
|
|
|
|
resource_access_hook: None,
|
|
|
|
|
|
}
|
|
|
|
|
|
}
|
|
|
|
|
|
|
|
|
|
|
|
/// Builder: attach a [`ResourceAccessHook`] that fires after every
|
|
|
|
|
|
/// authorised `_with_perms` read. Without it the service is silent —
|
|
|
|
|
|
/// existing behaviour for stub / test paths.
|
|
|
|
|
|
pub fn with_resource_access_hook(mut self, hook: Arc<dyn ResourceAccessHook>) -> Self {
|
|
|
|
|
|
self.resource_access_hook = Some(hook);
|
|
|
|
|
|
self
|
|
|
|
|
|
}
|
|
|
|
|
|
|
|
|
|
|
|
/// Fire the access hook if registered. Called from every `_with_perms`
|
|
|
|
|
|
/// read after the authZ + lookup has succeeded (never on failure
|
|
|
|
|
|
/// paths — denied reads must not surface in Recent).
|
|
|
|
|
|
///
|
|
|
|
|
|
/// `pub` because the WebDAV / NextCloud DAV handlers resolve files
|
|
|
|
|
|
/// by path and authorise via that resolver, not via the
|
|
|
|
|
|
/// `*_with_perms` service methods — they then serve content through
|
|
|
|
|
|
/// the no-perms `get_file_stream` / `get_file_range_stream`. Those
|
|
|
|
|
|
/// handlers must call this directly after their own authZ has
|
|
|
|
|
|
/// passed so cross-protocol downloads (NC desktop, davx5, native
|
|
|
|
|
|
/// `/webdav/`) also surface in Recent.
|
|
|
|
|
|
pub fn notify_file_accessed(&self, caller_id: Uuid, file_id: &str) {
|
|
|
|
|
|
if let Some(hook) = &self.resource_access_hook {
|
|
|
|
|
|
hook.on_file_accessed(caller_id, file_id);
|
2026-02-14 17:54:25 +01:00
|
|
|
|
}
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
// ── private helpers ──────────────────────────────────────────
|
|
|
|
|
|
|
2026-06-20 14:42:10 +02:00
|
|
|
|
/// Read a file's full content through the streaming API into a single
|
|
|
|
|
|
/// `Bytes` buffer. Working memory stays at one chunk while reading; the
|
|
|
|
|
|
/// returned buffer holds the whole (sub-threshold) file.
|
|
|
|
|
|
async fn read_full(
|
|
|
|
|
|
file_read: &FileBlobReadRepository,
|
|
|
|
|
|
id: &str,
|
|
|
|
|
|
capacity: usize,
|
|
|
|
|
|
) -> Result<Bytes, DomainError> {
|
|
|
|
|
|
let stream = file_read.get_file_stream(id).await?;
|
|
|
|
|
|
let mut stream = Pin::from(stream);
|
|
|
|
|
|
let mut buf = BytesMut::with_capacity(capacity);
|
|
|
|
|
|
while let Some(chunk) = stream.next().await {
|
|
|
|
|
|
buf.extend_from_slice(&chunk.map_err(|e| {
|
|
|
|
|
|
DomainError::internal_error("File", format!("Stream read error: {}", e))
|
|
|
|
|
|
})?);
|
|
|
|
|
|
}
|
|
|
|
|
|
Ok(buf.freeze())
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-20 22:56:00 +02:00
|
|
|
|
/// Helper: require the caller has `perm` on the given file id.
|
|
|
|
|
|
/// Fail-closed if no engine was injected (stub/test path).
|
|
|
|
|
|
async fn require_file(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
file_id: &str,
|
|
|
|
|
|
perm: Permission,
|
|
|
|
|
|
caller_id: Uuid,
|
|
|
|
|
|
) -> Result<(), DomainError> {
|
|
|
|
|
|
let authz = self.authz.as_ref().ok_or_else(|| {
|
|
|
|
|
|
DomainError::internal_error("FileRetrieval", "Authorization engine unavailable")
|
|
|
|
|
|
})?;
|
|
|
|
|
|
let uuid = Uuid::parse_str(file_id).map_err(|_| DomainError::not_found("File", file_id))?;
|
|
|
|
|
|
authz
|
|
|
|
|
|
.require(Subject::User(caller_id), perm, Resource::File(uuid))
|
|
|
|
|
|
.await
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
/// Engine check for a target folder. `None` is allowed (root namespace,
|
|
|
|
|
|
/// implicitly owned by the caller).
|
|
|
|
|
|
async fn require_target_folder_perm(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
folder_id: Option<&str>,
|
|
|
|
|
|
perm: Permission,
|
|
|
|
|
|
caller_id: Uuid,
|
|
|
|
|
|
) -> Result<(), DomainError> {
|
|
|
|
|
|
let Some(target) = folder_id else {
|
|
|
|
|
|
return Ok(());
|
|
|
|
|
|
};
|
|
|
|
|
|
let authz = self.authz.as_ref().ok_or_else(|| {
|
|
|
|
|
|
DomainError::internal_error("FileRetrieval", "Authorization engine unavailable")
|
|
|
|
|
|
})?;
|
|
|
|
|
|
let uuid = Uuid::parse_str(target).map_err(|_| DomainError::not_found("Folder", target))?;
|
|
|
|
|
|
authz
|
|
|
|
|
|
.require(Subject::User(caller_id), perm, Resource::Folder(uuid))
|
|
|
|
|
|
.await
|
|
|
|
|
|
}
|
2026-02-08 13:40:23 +01:00
|
|
|
|
|
|
|
|
|
|
/// Try to transcode image content to WebP and return transcoded variant.
|
|
|
|
|
|
async fn try_transcode(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
|
|
|
|
|
content: &Bytes,
|
|
|
|
|
|
mime: &str,
|
|
|
|
|
|
file_size: u64,
|
|
|
|
|
|
accept_webp: bool,
|
2026-03-03 15:55:15 +00:00
|
|
|
|
) -> Option<(Bytes, Arc<str>)> {
|
2026-02-08 13:40:23 +01:00
|
|
|
|
if !accept_webp {
|
|
|
|
|
|
return None;
|
|
|
|
|
|
}
|
|
|
|
|
|
let transcode = self.transcode.as_ref()?;
|
2026-03-03 15:36:42 +00:00
|
|
|
|
if !ImageTranscodeService::should_transcode(mime, file_size) {
|
2026-02-08 13:40:23 +01:00
|
|
|
|
return None;
|
|
|
|
|
|
}
|
|
|
|
|
|
let format = OutputFormat::WebP;
|
2026-02-25 10:28:34 +01:00
|
|
|
|
match transcode
|
|
|
|
|
|
.get_transcoded(id, content.clone(), mime, format)
|
|
|
|
|
|
.await
|
|
|
|
|
|
{
|
2026-02-08 13:40:23 +01:00
|
|
|
|
Ok((transcoded, webp_mime, true)) => {
|
|
|
|
|
|
debug!(
|
|
|
|
|
|
"🖼️ WebP transcode: {} -> {} bytes ({:.0}% smaller)",
|
|
|
|
|
|
content.len(),
|
|
|
|
|
|
transcoded.len(),
|
|
|
|
|
|
(1.0 - transcoded.len() as f64 / content.len().max(1) as f64) * 100.0
|
|
|
|
|
|
);
|
2026-03-03 15:55:15 +00:00
|
|
|
|
Some((transcoded, Arc::from(&*webp_mime)))
|
2026-02-08 13:40:23 +01:00
|
|
|
|
}
|
|
|
|
|
|
_ => None,
|
|
|
|
|
|
}
|
2025-03-19 19:52:12 +01:00
|
|
|
|
}
|
2026-02-08 13:40:23 +01:00
|
|
|
|
|
2026-02-15 17:53:25 +01:00
|
|
|
|
/// Core multi-tier download logic shared by `get_file_optimized` and
|
|
|
|
|
|
/// `get_file_optimized_preloaded`.
|
|
|
|
|
|
async fn optimized_inner(
|
2026-02-08 13:40:23 +01:00
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
2026-02-15 17:53:25 +01:00
|
|
|
|
dto: FileDto,
|
2026-02-08 13:40:23 +01:00
|
|
|
|
accept_webp: bool,
|
|
|
|
|
|
prefer_original: bool,
|
|
|
|
|
|
) -> Result<(FileDto, OptimizedFileContent), DomainError> {
|
|
|
|
|
|
let mime_type = dto.mime_type.clone();
|
|
|
|
|
|
let file_size = dto.size;
|
|
|
|
|
|
let file_name = dto.name.clone();
|
2026-06-15 11:43:15 +00:00
|
|
|
|
// The content cache is content-addressed: keyed by the blob hash, not
|
|
|
|
|
|
// the file id. Identical content deduplicated to one blob on disk is
|
|
|
|
|
|
// then cached ONCE in RAM and shared by every file/user that references
|
|
|
|
|
|
// it — the cache benefits from dedup, not just the disk. Immutable by
|
|
|
|
|
|
// construction, so entries never go stale (no invalidation needed). A
|
|
|
|
|
|
// stub DTO without a hash disables caching for that request rather than
|
|
|
|
|
|
// colliding every hash-less file on the key "".
|
|
|
|
|
|
let cache_key = dto.content_hash.clone();
|
|
|
|
|
|
let cacheable = !cache_key.is_empty();
|
2026-02-08 13:40:23 +01:00
|
|
|
|
let do_transcode = accept_webp && !prefer_original;
|
|
|
|
|
|
|
|
|
|
|
|
// ── Tier 1: Hot cache + transcode (<10 MB) ──────────
|
|
|
|
|
|
if file_size < CACHE_THRESHOLD {
|
2026-06-20 14:42:10 +02:00
|
|
|
|
// Fetch the raw blob bytes. When cacheable, `get_or_load` serves
|
|
|
|
|
|
// from the content cache on a hit and, on a miss, coalesces every
|
|
|
|
|
|
// concurrent request for the same blob hash into a SINGLE disk read
|
|
|
|
|
|
// (single-flight) — no thundering herd under load. Hash-less stub
|
|
|
|
|
|
// DTOs are uncacheable and stream straight from disk.
|
|
|
|
|
|
let content_bytes = if cacheable && let Some(cache) = &self.content_cache {
|
2026-06-15 11:43:15 +00:00
|
|
|
|
let etag: Arc<str> = format!("\"{}\"", cache_key).into();
|
2026-03-03 15:55:15 +00:00
|
|
|
|
let ct: Arc<str> = mime_type.clone();
|
2026-06-20 14:42:10 +02:00
|
|
|
|
let file_read = Arc::clone(&self.file_read);
|
|
|
|
|
|
let id_owned = id.to_string();
|
|
|
|
|
|
let cap = file_size as usize;
|
|
|
|
|
|
let (bytes, _etag, _ct) = cache
|
|
|
|
|
|
.get_or_load(cache_key.clone(), etag, ct, async move {
|
|
|
|
|
|
debug!("💾 TIER 1 Cache MISS: {} – loading from disk", id_owned);
|
|
|
|
|
|
Self::read_full(&file_read, &id_owned, cap).await
|
|
|
|
|
|
})
|
|
|
|
|
|
.await?;
|
|
|
|
|
|
bytes
|
|
|
|
|
|
} else {
|
|
|
|
|
|
debug!(
|
|
|
|
|
|
"💾 TIER 1 (uncacheable): {} – streaming from disk",
|
|
|
|
|
|
file_name
|
|
|
|
|
|
);
|
|
|
|
|
|
Self::read_full(&self.file_read, id, file_size as usize).await?
|
|
|
|
|
|
};
|
2026-02-08 13:40:23 +01:00
|
|
|
|
|
2026-02-14 01:26:02 +01:00
|
|
|
|
if do_transcode
|
2026-02-14 01:29:34 +01:00
|
|
|
|
&& let Some((t, m)) = self
|
|
|
|
|
|
.try_transcode(id, &content_bytes, &mime_type, file_size, true)
|
|
|
|
|
|
.await
|
|
|
|
|
|
{
|
|
|
|
|
|
return Ok((
|
|
|
|
|
|
dto,
|
|
|
|
|
|
OptimizedFileContent::Bytes {
|
2026-02-08 13:40:23 +01:00
|
|
|
|
data: t,
|
|
|
|
|
|
mime_type: m,
|
|
|
|
|
|
was_transcoded: true,
|
2026-02-14 01:29:34 +01:00
|
|
|
|
},
|
|
|
|
|
|
));
|
|
|
|
|
|
}
|
|
|
|
|
|
return Ok((
|
|
|
|
|
|
dto,
|
|
|
|
|
|
OptimizedFileContent::Bytes {
|
|
|
|
|
|
data: content_bytes,
|
|
|
|
|
|
mime_type: mime_type.clone(),
|
|
|
|
|
|
was_transcoded: false,
|
|
|
|
|
|
},
|
|
|
|
|
|
));
|
2026-02-08 13:40:23 +01:00
|
|
|
|
}
|
|
|
|
|
|
|
2026-02-15 17:53:25 +01:00
|
|
|
|
// ── Tier 2 + 3: Streaming (≥10 MB) ──────────────────
|
2026-02-14 01:29:34 +01:00
|
|
|
|
info!(
|
2026-02-15 17:53:25 +01:00
|
|
|
|
"📡 TIER 2 STREAMING: {} ({} MB)",
|
2026-02-14 01:29:34 +01:00
|
|
|
|
file_name,
|
|
|
|
|
|
file_size / (1024 * 1024)
|
|
|
|
|
|
);
|
2026-02-23 00:51:46 +01:00
|
|
|
|
let stream = self.file_read.get_file_stream(id).await?;
|
|
|
|
|
|
Ok((dto, OptimizedFileContent::Stream(Box::into_pin(stream))))
|
2026-02-08 13:40:23 +01:00
|
|
|
|
}
|
2026-06-19 11:27:11 +00:00
|
|
|
|
|
|
|
|
|
|
/// Batch counterpart of [`FileRetrievalUseCase::get_file`]: resolve many
|
|
|
|
|
|
/// file ids in ONE query instead of one per id. Like `get_file` it
|
|
|
|
|
|
/// performs no per-file authorization — both current callers (ACL grant
|
|
|
|
|
|
/// listing, NextCloud favorites REPORT) resolve ids already vetted by the
|
|
|
|
|
|
/// authorization engine or the favorites table. Missing or trashed ids are
|
|
|
|
|
|
/// absent from the result; callers re-associate by `id`.
|
|
|
|
|
|
pub async fn get_files_by_ids(&self, ids: &[String]) -> Result<Vec<FileDto>, DomainError> {
|
|
|
|
|
|
let files = self.file_read.get_files_by_ids(ids).await?;
|
|
|
|
|
|
Ok(files.into_iter().map(FileDto::from).collect())
|
|
|
|
|
|
}
|
2026-02-15 17:53:25 +01:00
|
|
|
|
}
|
|
|
|
|
|
|
|
|
|
|
|
impl FileRetrievalUseCase for FileRetrievalService {
|
|
|
|
|
|
async fn get_file(&self, id: &str) -> Result<FileDto, DomainError> {
|
|
|
|
|
|
let file = self.file_read.get_file(id).await?;
|
|
|
|
|
|
Ok(FileDto::from(file))
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
async fn get_file_with_perms(&self, id: &str, caller_id: Uuid) -> Result<FileDto, DomainError> {
|
2026-05-20 22:56:00 +02:00
|
|
|
|
self.require_file(id, Permission::Read, caller_id).await?;
|
|
|
|
|
|
let file = self.file_read.get_file(id).await?;
|
2026-06-26 01:01:44 +02:00
|
|
|
|
// After authZ + lookup succeed: this caller has just inspected the
|
|
|
|
|
|
// file. Recent listing observes via the hook. The throttle in the
|
|
|
|
|
|
// recording impl coalesces repeat metadata fetches against the same
|
|
|
|
|
|
// file (file viewer poll, browse-then-download pattern).
|
|
|
|
|
|
self.notify_file_accessed(caller_id, id);
|
2026-03-04 17:18:39 +01:00
|
|
|
|
Ok(FileDto::from(file))
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-22 00:38:44 +02:00
|
|
|
|
async fn get_file_or_trashed_with_perms(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
|
|
|
|
|
caller_id: Uuid,
|
|
|
|
|
|
) -> Result<FileDto, DomainError> {
|
|
|
|
|
|
self.require_file(id, Permission::Read, caller_id).await?;
|
|
|
|
|
|
let file = self.file_read.get_file_or_trashed(id).await?;
|
|
|
|
|
|
Ok(FileDto::from(file))
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
// FIXME no authorisation at all
|
2026-06-19 07:49:33 +02:00
|
|
|
|
async fn get_file_by_path(&self, path: &str, drive_id: Uuid) -> Result<FileDto, DomainError> {
|
2026-02-15 17:53:25 +01:00
|
|
|
|
// Direct SQL lookup — O(folder_depth) queries instead of O(total_files)
|
2026-05-20 22:56:00 +02:00
|
|
|
|
// NOTE: This method does NOT perform any authorization check. Callers
|
|
|
|
|
|
// that surface its result to a user-driven request MUST resolve the
|
|
|
|
|
|
// file via get_file_owned afterwards, or call authz.require directly.
|
|
|
|
|
|
// (Tracked in the audit punch-list under "path-based lookups".)
|
2026-06-19 07:49:33 +02:00
|
|
|
|
// `drive_id` scope axis prevents cross-drive resolution — without
|
|
|
|
|
|
// it, `find_file_by_path` would return a non-deterministic row
|
|
|
|
|
|
// when the same path exists in multiple drives.
|
|
|
|
|
|
if let Some(file) = self.file_read.find_file_by_path(path, drive_id).await? {
|
2026-02-15 17:53:25 +01:00
|
|
|
|
return Ok(FileDto::from(file));
|
|
|
|
|
|
}
|
|
|
|
|
|
|
|
|
|
|
|
Err(DomainError::not_found(
|
|
|
|
|
|
"File",
|
|
|
|
|
|
format!("not found at path: {}", path),
|
|
|
|
|
|
))
|
|
|
|
|
|
}
|
|
|
|
|
|
|
|
|
|
|
|
async fn list_files(&self, folder_id: Option<&str>) -> Result<Vec<FileDto>, DomainError> {
|
|
|
|
|
|
let files = self.file_read.list_files(folder_id).await?;
|
|
|
|
|
|
Ok(files.into_iter().map(FileDto::from).collect())
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
async fn list_files_with_perms(
|
2026-03-05 10:30:39 +01:00
|
|
|
|
&self,
|
|
|
|
|
|
folder_id: Option<&str>,
|
2026-03-07 14:59:32 +01:00
|
|
|
|
owner_id: Uuid,
|
2026-03-05 10:30:39 +01:00
|
|
|
|
) -> Result<Vec<FileDto>, DomainError> {
|
2026-05-21 11:07:04 +02:00
|
|
|
|
if folder_id.is_some() {
|
|
|
|
|
|
// folder id is defined, check permissions
|
|
|
|
|
|
self.require_target_folder_perm(folder_id, Permission::Read, owner_id)
|
|
|
|
|
|
.await?;
|
|
|
|
|
|
self.list_files(folder_id).await
|
|
|
|
|
|
} else {
|
|
|
|
|
|
// no folder id, get owners's files' root
|
|
|
|
|
|
let files = self
|
|
|
|
|
|
.file_read
|
|
|
|
|
|
.list_files_for_owner(folder_id, owner_id)
|
|
|
|
|
|
.await?;
|
|
|
|
|
|
Ok(files.into_iter().map(FileDto::from).collect())
|
|
|
|
|
|
}
|
2026-03-05 10:30:39 +01:00
|
|
|
|
}
|
|
|
|
|
|
|
2026-02-15 17:53:25 +01:00
|
|
|
|
async fn get_file_stream(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
|
|
|
|
|
) -> Result<Box<dyn Stream<Item = Result<Bytes, std::io::Error>> + Send>, DomainError> {
|
|
|
|
|
|
self.file_read.get_file_stream(id).await
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
async fn get_file_stream_with_perms(
|
2026-03-05 10:30:39 +01:00
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
2026-03-07 14:59:32 +01:00
|
|
|
|
caller_id: Uuid,
|
2026-03-05 10:30:39 +01:00
|
|
|
|
) -> Result<Box<dyn Stream<Item = Result<Bytes, std::io::Error>> + Send>, DomainError> {
|
2026-05-20 22:56:00 +02:00
|
|
|
|
self.require_file(id, Permission::Read, caller_id).await?;
|
2026-06-26 01:01:44 +02:00
|
|
|
|
self.notify_file_accessed(caller_id, id);
|
2026-03-05 10:30:39 +01:00
|
|
|
|
self.file_read.get_file_stream(id).await
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-02-15 17:53:25 +01:00
|
|
|
|
/// Multi-tier optimized download.
|
|
|
|
|
|
async fn get_file_optimized(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
|
|
|
|
|
accept_webp: bool,
|
|
|
|
|
|
prefer_original: bool,
|
|
|
|
|
|
) -> Result<(FileDto, OptimizedFileContent), DomainError> {
|
|
|
|
|
|
let file = self.file_read.get_file(id).await?;
|
|
|
|
|
|
let dto = FileDto::from(file);
|
|
|
|
|
|
self.optimized_inner(id, dto, accept_webp, prefer_original)
|
|
|
|
|
|
.await
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
async fn get_file_optimized_with_perms(
|
2026-03-04 17:18:39 +01:00
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
2026-03-07 14:59:32 +01:00
|
|
|
|
caller_id: Uuid,
|
2026-03-04 17:18:39 +01:00
|
|
|
|
accept_webp: bool,
|
|
|
|
|
|
prefer_original: bool,
|
|
|
|
|
|
) -> Result<(FileDto, OptimizedFileContent), DomainError> {
|
2026-05-20 22:56:00 +02:00
|
|
|
|
self.require_file(id, Permission::Read, caller_id).await?;
|
|
|
|
|
|
let file = self.file_read.get_file(id).await?;
|
2026-03-04 17:18:39 +01:00
|
|
|
|
let dto = FileDto::from(file);
|
2026-06-26 01:01:44 +02:00
|
|
|
|
self.notify_file_accessed(caller_id, id);
|
2026-03-04 17:18:39 +01:00
|
|
|
|
self.optimized_inner(id, dto, accept_webp, prefer_original)
|
|
|
|
|
|
.await
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-02-15 17:53:25 +01:00
|
|
|
|
/// Like `get_file_optimized` but skips the metadata re-fetch.
|
|
|
|
|
|
async fn get_file_optimized_preloaded(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
|
|
|
|
|
file_dto: FileDto,
|
|
|
|
|
|
accept_webp: bool,
|
|
|
|
|
|
prefer_original: bool,
|
|
|
|
|
|
) -> Result<(FileDto, OptimizedFileContent), DomainError> {
|
|
|
|
|
|
self.optimized_inner(id, file_dto, accept_webp, prefer_original)
|
|
|
|
|
|
.await
|
|
|
|
|
|
}
|
2026-02-08 13:40:23 +01:00
|
|
|
|
|
|
|
|
|
|
/// Range-based streaming for HTTP Range Requests.
|
|
|
|
|
|
async fn get_file_range_stream(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
|
|
|
|
|
start: u64,
|
|
|
|
|
|
end: Option<u64>,
|
|
|
|
|
|
) -> Result<Box<dyn Stream<Item = Result<Bytes, std::io::Error>> + Send>, DomainError> {
|
|
|
|
|
|
self.file_read.get_file_range_stream(id, start, end).await
|
2025-03-19 19:52:12 +01:00
|
|
|
|
}
|
2026-02-24 12:18:38 +01:00
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
async fn get_file_range_stream_with_perms(
|
2026-03-04 17:18:39 +01:00
|
|
|
|
&self,
|
|
|
|
|
|
id: &str,
|
2026-03-07 14:59:32 +01:00
|
|
|
|
caller_id: Uuid,
|
2026-03-04 17:18:39 +01:00
|
|
|
|
start: u64,
|
|
|
|
|
|
end: Option<u64>,
|
|
|
|
|
|
) -> Result<Box<dyn Stream<Item = Result<Bytes, std::io::Error>> + Send>, DomainError> {
|
2026-05-20 22:56:00 +02:00
|
|
|
|
self.require_file(id, Permission::Read, caller_id).await?;
|
2026-06-26 01:01:44 +02:00
|
|
|
|
// Range requests are bursty (video seeks, NC chunked downloads) —
|
|
|
|
|
|
// the recording hook's per-(caller, file) throttle absorbs the
|
|
|
|
|
|
// storm so one watched video lands as one Recent row, not 1000.
|
|
|
|
|
|
self.notify_file_accessed(caller_id, id);
|
2026-03-04 17:18:39 +01:00
|
|
|
|
self.file_read.get_file_range_stream(id, start, end).await
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
// TODO: check: no permission check
|
2026-02-26 00:07:10 +01:00
|
|
|
|
async fn stream_files_in_subtree(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
folder_id: &str,
|
|
|
|
|
|
) -> Result<Pin<Box<dyn Stream<Item = Result<FileDto, DomainError>> + Send>>, DomainError> {
|
|
|
|
|
|
let inner = self.file_read.stream_files_in_subtree(folder_id).await?;
|
|
|
|
|
|
let mapped = inner.map(|r| r.map(FileDto::from));
|
|
|
|
|
|
Ok(Box::pin(mapped))
|
2026-02-24 12:18:38 +01:00
|
|
|
|
}
|
2026-02-24 15:11:56 +01:00
|
|
|
|
|
|
|
|
|
|
async fn list_files_batch(
|
|
|
|
|
|
&self,
|
|
|
|
|
|
folder_id: Option<&str>,
|
|
|
|
|
|
offset: i64,
|
|
|
|
|
|
limit: i64,
|
|
|
|
|
|
) -> Result<Vec<FileDto>, DomainError> {
|
2026-02-25 10:28:34 +01:00
|
|
|
|
let files = self
|
|
|
|
|
|
.file_read
|
|
|
|
|
|
.list_files_batch(folder_id, offset, limit)
|
|
|
|
|
|
.await?;
|
2026-02-24 15:11:56 +01:00
|
|
|
|
Ok(files.into_iter().map(FileDto::from).collect())
|
|
|
|
|
|
}
|
2026-03-05 10:30:39 +01:00
|
|
|
|
|
2026-05-21 11:07:04 +02:00
|
|
|
|
async fn list_files_batch_with_perms(
|
2026-03-05 10:30:39 +01:00
|
|
|
|
&self,
|
|
|
|
|
|
folder_id: Option<&str>,
|
2026-03-07 14:59:32 +01:00
|
|
|
|
owner_id: Uuid,
|
2026-03-05 10:30:39 +01:00
|
|
|
|
offset: i64,
|
|
|
|
|
|
limit: i64,
|
|
|
|
|
|
) -> Result<Vec<FileDto>, DomainError> {
|
2026-05-21 11:07:04 +02:00
|
|
|
|
if folder_id.is_some() {
|
|
|
|
|
|
// folder id is defined, check permissions
|
|
|
|
|
|
self.require_target_folder_perm(folder_id, Permission::Read, owner_id)
|
|
|
|
|
|
.await?;
|
|
|
|
|
|
let files = self
|
|
|
|
|
|
.file_read
|
|
|
|
|
|
.list_files_batch(folder_id, offset, limit)
|
|
|
|
|
|
.await?;
|
|
|
|
|
|
return Ok(files.into_iter().map(FileDto::from).collect());
|
|
|
|
|
|
}
|
|
|
|
|
|
|
2026-03-05 10:30:39 +01:00
|
|
|
|
let files = self
|
|
|
|
|
|
.file_read
|
|
|
|
|
|
.list_files_batch_for_owner(folder_id, owner_id, offset, limit)
|
|
|
|
|
|
.await?;
|
|
|
|
|
|
Ok(files.into_iter().map(FileDto::from).collect())
|
|
|
|
|
|
}
|
2026-02-14 01:29:34 +01:00
|
|
|
|
}
|