{ inputs, ... }: let # Load workspace from uv.lock workspace = inputs.uv2nix.lib.workspace.loadWorkspace { workspaceRoot = ../.; }; mkPythonSet = { pkgs, lib, self', members }: let inherit (pkgs.stdenv.hostPlatform) isLinux isDarwin isx86_64; inherit (pkgs.config) cudaSupport; inherit (pkgs) cudaPackages; cuda13Support = cudaSupport && cudaPackages.cudaMajorVersion == "13"; libmlx_source = if cuda13Support then "mlx-cuda-13" else if cudaSupport then "mlx-cuda-12" else "mlx-cpu"; python = pkgs.python313; cudaLibs = with cudaPackages; [ cuda_cudart cuda_cccl cuda_cupti cuda_nvrtc cuda_nvtx cudnn libcufile libcublas libcufft libcurand libcusolver libcusparse libcusparse_lt libnvjitlink libnvshmem nccl ]; exoOverlay = final: prev: { # Replace workspace exo_pyo3_bindings with Nix-built wheel. # Preserve passthru so mkVirtualEnv can resolve dependency groups. # Copy .pyi stub + py.typed marker so basedpyright can find the types. exo-pyo3-bindings = pkgs.stdenv.mkDerivation { pname = "exo-pyo3-bindings"; version = "0.1.0"; src = self'.packages.exo_pyo3_bindings; # Install from pre-built wheel nativeBuildInputs = [ final.pyprojectWheelHook ]; dontStrip = true; passthru = prev.exo-pyo3-bindings.passthru or { }; postInstall = '' local siteDir=$out/${final.python.sitePackages}/exo_pyo3_bindings cp ${inputs.self}/rust/exo_pyo3_bindings/exo_pyo3_bindings.pyi $siteDir/ touch $siteDir/py.typed ''; }; }; buildSystemsOverlay = final: prev: lib.optionalAttrs isDarwin { mlx = prev.mlx.overrideAttrs (old: let # Static dependencies included directly during compilation gguf-tools = pkgs.fetchFromGitHub { owner = "antirez"; repo = "gguf-tools"; rev = "8fa6eb65236618e28fd7710a0fba565f7faa1848"; hash = "sha256-15FvyPOFqTOr5vdWQoPnZz+mYH919++EtghjozDlnSA="; }; metal_cpp = pkgs.fetchzip { url = "https://developer.apple.com/metal/cpp/files/metal-cpp_26.zip"; hash = "sha256-7n2eI2lw/S+Us6l7YPAATKwcIbRRpaQ8VmES7S8ZjY8="; }; nanobind = pkgs.fetchFromGitHub { owner = "wjakob"; repo = "nanobind"; rev = "v2.10.2"; hash = "sha256-io44YhN+VpfHFWyvvLWSanRgbzA0whK8WlDNRi3hahU="; fetchSubmodules = true; }; in { nativeBuildInputs = (old.nativeBuildInputs or [ ]) ++ [ pkgs.cmake self'.packages.metal-toolchain ]; # TODO: non-sdk_26 support buildInputs = (old.buildInputs or [ ]) ++ [ gguf-tools pkgs.fmt pkgs.nlohmann_json pkgs.apple-sdk_26 ]; patches = [ (pkgs.replaceVars ../nix/darwin-build-fixes.patch { sdkVersion = pkgs.apple-sdk_26.version; inherit (self'.packages.metal-toolchain) metalVersion; }) ]; postPatch = '' substituteInPlace mlx/backend/cpu/jit_compiler.cpp \ --replace-fail "g++" "${lib.getExe' pkgs.stdenv.cc "c++"}" ''; DEV_RELEASE = 1; CMAKE_ARGS = toString ([ (lib.cmakeBool "USE_SYSTEM_FMT" true) (lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_GGUFLIB" "${gguf-tools}") (lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_JSON" "${pkgs.nlohmann_json.src}") (lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_NANOBIND" "${nanobind}") (lib.cmakeBool "FETCHCONTENT_FULLY_DISCONNECTED" true) (lib.cmakeBool "MLX_BUILD_CPU" true) (lib.cmakeBool "MLX_BUILD_METAL" true) (lib.cmakeOptionType "string" "CMAKE_INSTALL_LIBDIR" "lib") (lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_METAL_CPP" "${metal_cpp}") (lib.cmakeOptionType "string" "CMAKE_OSX_DEPLOYMENT_TARGET" "${pkgs.apple-sdk_26.version}") (lib.cmakeOptionType "filepath" "CMAKE_OSX_SYSROOT" "${pkgs.apple-sdk_26.passthru.sdkroot}") ] ++ lib.optionals (isDarwin && isx86_64) [ (lib.cmakeBool "MLX_ENABLE_X64_MAC" true) ]); SDKROOT = pkgs.apple-sdk_26.passthru.sdkroot; MACOSX_DEPLOYMENT_TARGET = pkgs.apple-sdk_26.version; }); } // lib.optionalAttrs isLinux { mlx = prev.mlx.overrideAttrs (old: { buildInputs = old.buildInputs ++ lib.optionals cudaSupport cudaLibs; autoPatchelfIgnoreMissingDeps = lib.optionals cudaSupport [ "libcuda.so.1" ]; postInstall = '' cp -r "${final.${libmlx_source}}/${final.python.sitePackages}/mlx" "$out/${final.python.sitePackages}/mlx/" ''; }); } // lib.optionalAttrs cudaSupport { "${libmlx_source}" = prev."${libmlx_source}".overrideAttrs (old: { buildInputs = old.buildInputs ++ cudaLibs; autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ]; }); nvidia-cufile = prev.nvidia-cufile.overrideAttrs (old: { buildInputs = old.buildInputs ++ [ pkgs.rdma-core ]; autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ]; }); nvidia-cusolver = prev.nvidia-cusolver.overrideAttrs (old: { buildInputs = old.buildInputs ++ cudaLibs; autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ]; }); nvidia-nvshmem-cu13 = prev.nvidia-nvshmem-cu13.overrideAttrs (old: { buildInputs = old.buildInputs ++ [ pkgs.rdma-core pkgs.pmix pkgs.libfabric pkgs.ucx pkgs.openmpi ]; autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ]; }); nvidia-cusparse = prev.nvidia-cusparse.overrideAttrs (old: { buildInputs = old.buildInputs ++ [ cudaLibs ]; autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ]; }); torch = prev.torch.overrideAttrs (old: { buildInputs = old.buildInputs ++ cudaLibs; autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ]; }); }; pyprojectOverlay = workspace.mkPyprojectOverlay { sourcePreference = "wheel"; dependencies = members; }; editableOverlay = workspace.mkEditablePyprojectOverlay { # Use environment variable pointing to editable root directory root = "$REPO_ROOT"; members = [ "exo" "exo-bench" ]; }; pythonSet = (pkgs.callPackage inputs.pyproject-nix.build.packages { inherit python; }).overrideScope ( lib.composeManyExtensions [ inputs.pyproject-build-systems.overlays.default pyprojectOverlay exoOverlay buildSystemsOverlay ] ); venv = name: (pythonSet.mkVirtualEnv "${name}-env" members).overrideAttrs (_: { venvSkip = [ "lib/python${python.pythonVersion}/site-packages/mlx/share/cmake/*" ]; }); mkApp = cmd: name: pkgs.writeShellApplication { inherit name; runtimeEnv = { EXO_DASHBOARD_DIR = self'.packages.dashboard; EXO_RESOURCES_DIR = inputs.self + /resources; }; runtimeInputs = [ # mlx and mlx-cuda ship clashing cmake files - we dont need them at runtime anyway (venv name) ] ++ lib.optionals isDarwin [ pkgs.macmon ]; text = "exec " + lib.optionalString cudaSupport "${lib.getExe pkgs.nix-gl-host} " + cmd; }; in { inherit venv; editablePythonSet = pythonSet.overrideScope editableOverlay; mkPythonScript = path: mkApp ''python ${path} "$@"''; mkExo = mkApp ''exo "$@"''; }; in { perSystem = { self', pkgs, unfreePkgs, lib, ... }: let inherit (pkgs.stdenv.hostPlatform) isLinux; inherit (mkPythonSet { inherit self' pkgs lib; members = { exo = [ "cpu" ]; }; }) editablePythonSet mkExo; # Virtual environment with dev dependencies for testing testVenv = (mkPythonSet { inherit self' pkgs lib; members = { exo = [ "dev" "cpu" ]; # Include pytest, pytest-asyncio, pytest-env }; }).venv "exo-test"; mkBenchScript = (mkPythonSet { inherit self' pkgs lib; members = { exo = [ "cpu" ]; exo-bench = [ ]; # Include pytest, pytest-asyncio, pytest-env }; }).mkPythonScript; mkSimplePythonScript = name: path: pkgs.writeShellApplication { inherit name; runtimeInputs = [ pkgs.python313 ]; text = ''exec python ${path} "$@"''; }; in { packages = { exo = mkExo "exo"; editableVenv = editablePythonSet.mkVirtualEnv "exo-dev-env" { exo = [ "dev" ]; }; # for running tests in ci exo-test-env = testVenv; exo-bench = mkBenchScript "exo-bench" (inputs.self + /bench/exo_bench.py); exo-eval = mkBenchScript "exo-eval" (inputs.self + /bench/exo_eval.py); exo-eval-tool-calls = mkBenchScript "exo-eval-tool-calls" (inputs.self + /bench/eval_tool_calls.py); # used by ./tests/run_exo_on.sh exo-get-all-models-on-cluster = mkSimplePythonScript "exo-get-all-models-on-cluster" (inputs.self + /tests/get_all_models_on_cluster.py); } // lib.optionalAttrs isLinux { exo-cuda-12 = (mkPythonSet { inherit self' lib; inherit (unfreePkgs.pkgsCuda.cudaPackages_12) pkgs; members = { exo = [ "cuda12" ]; }; }).mkExo "exo-cuda-12"; exo-cuda-13 = (mkPythonSet { inherit self' lib; inherit (unfreePkgs.pkgsCuda.cudaPackages_13) pkgs; members = { exo = [ "cuda13" ]; }; }).mkExo "exo-cuda-13"; }; checks = { lint = pkgs.runCommand "ruff-lint" { } '' export RUFF_CACHE_DIR="$TMPDIR/ruff-cache" ${pkgs.ruff}/bin/ruff check ${inputs.self} touch $out ''; typecheck = pkgs.runCommand "typecheck" { nativeBuildInputs = [ testVenv ]; } '' cd ${inputs.self} basedpyright touch $out ''; }; }; }