From a8a8d849238afe89226adc18a4abbf372a7c3087 Mon Sep 17 00:00:00 2001 From: Claude Date: Fri, 12 Jun 2026 11:36:46 +0000 Subject: [PATCH] =?UTF-8?q?ASR=20:=20installation=20des=20d=C3=A9codeurs?= =?UTF-8?q?=20audio=20(soundfile,=20librosa)=20au=20d=C3=A9marrage=20de=20?= =?UTF-8?q?vLLM?= MIME-Version: 1.0 Content-Type: text/plain; charset=UTF-8 Content-Transfer-Encoding: 8bit L'image officielle vllm/vllm-openai ne contient pas les extras audio, ce qui faisait échouer /v1/audio/transcriptions en 400 'Invalid or unsupported audio file'. https://claude.ai/code/session_01YHMp3EKzr4s6o8w1ygxuUe --- docker-compose.unraid.yml | 6 +++++- docker-compose.yml | 6 +++++- 2 files changed, 10 insertions(+), 2 deletions(-) diff --git a/docker-compose.unraid.yml b/docker-compose.unraid.yml index e540cc2..95b5cbb 100644 --- a/docker-compose.unraid.yml +++ b/docker-compose.unraid.yml @@ -38,8 +38,12 @@ services: image: vllm/vllm-openai:latest container_name: liveflow-asr restart: unless-stopped + # L'image officielle vLLM n'embarque pas les décodeurs audio (soundfile, + # librosa) : on les installe au démarrage avant de lancer le serveur. + entrypoint: ["/bin/bash", "-c"] command: >- - --model Qwen/Qwen3-ASR-1.7B + pip install -q soundfile librosa && + exec vllm serve Qwen/Qwen3-ASR-1.7B --max-model-len 8192 --gpu-memory-utilization 0.70 runtime: nvidia diff --git a/docker-compose.yml b/docker-compose.yml index 4caa4b4..9743936 100644 --- a/docker-compose.yml +++ b/docker-compose.yml @@ -21,8 +21,12 @@ services: asr: image: vllm/vllm-openai:latest restart: unless-stopped + # L'image officielle vLLM n'embarque pas les décodeurs audio (soundfile, + # librosa) : on les installe au démarrage avant de lancer le serveur. + entrypoint: ["/bin/bash", "-c"] command: >- - --model Qwen/Qwen3-ASR-1.7B + pip install -q soundfile librosa && + exec vllm serve Qwen/Qwen3-ASR-1.7B --max-model-len 8192 --gpu-memory-utilization 0.70 environment: