# Loki — fork conteneurisé d'AJEAN (https://github.com/nathaninline/ajean, MIT) # Un seul service : le conteneur embarque le moteur llama.cpp (CUDA) ET l'UI. # # cp .env.example .env # ajuste si besoin # docker compose up --build # build local en quelques minutes (llama-server # # vient précompilé de l'image officielle llama.cpp) # # Interface : http://localhost:8090 # Pré-requis GPU : NVIDIA Container Toolkit sur l'hôte. services: loki: build: context: . # Pour épingler la version du moteur ou passer en CPU/Vulkan : # args: { LLAMACPP_IMAGE: "ghcr.io/ggml-org/llama.cpp:server-cuda-b10423" } container_name: loki ports: - "${LOKI_WEB_PORT:-8090}:${LOKI_WEB_PORT:-8090}" environment: - LOKI_WEB_PORT=${LOKI_WEB_PORT:-8090} # Semé UNE fois au premier démarrage (modifiable ensuite dans l'UI) : - LOKI_MODEL=${LOKI_MODEL:-} - LOKI_CTX=${LOKI_CTX:-} - LOKI_NGL=${LOKI_NGL:-} volumes: # /data = LOKI_HOME : config, base bbolt, mémoire, workspace, modèles téléchargés - ./data:/data # /models : dossier(s) de modèles GGUF supplémentaires (LOKI_MODEL_DIRS) - ./models:/models deploy: resources: reservations: devices: - driver: nvidia count: all capabilities: ["gpu"] restart: unless-stopped