mirror of
https://github.com/exo-explore/exo.git
synced 2026-04-18 04:52:40 -04:00
254 lines
11 KiB
Nix
254 lines
11 KiB
Nix
{ inputs, ... }:
|
|
let
|
|
# Load workspace from uv.lock
|
|
workspace = inputs.uv2nix.lib.workspace.loadWorkspace {
|
|
workspaceRoot = ../.;
|
|
};
|
|
|
|
mkPythonSet = { pkgs, lib, self' }:
|
|
let
|
|
inherit (pkgs.stdenv.hostPlatform) isLinux isDarwin isx86_64;
|
|
inherit (pkgs.config) cudaSupport;
|
|
inherit (pkgs) cudaPackages;
|
|
cuda13Support = cudaSupport && cudaPackages.cudaMajorVersion == "13";
|
|
libmlx_source = if cuda13Support then "mlx-cuda-13" else if cudaSupport then "mlx-cuda-12" else "mlx-cpu";
|
|
uv_extra = if cuda13Support then "cuda13" else if cudaSupport then "cuda12" else "cpu";
|
|
python = pkgs.python313;
|
|
cudaLibs = with cudaPackages; [
|
|
cuda_cudart
|
|
cuda_cccl
|
|
cuda_cupti
|
|
cuda_nvrtc
|
|
cuda_nvtx
|
|
cudnn
|
|
libcufile
|
|
libcublas
|
|
libcufft
|
|
libcurand
|
|
libcusolver
|
|
libcusparse
|
|
libcusparse_lt
|
|
libnvjitlink
|
|
libnvshmem
|
|
nccl
|
|
];
|
|
exoOverlay = final: prev: {
|
|
# Replace workspace exo_pyo3_bindings with Nix-built wheel.
|
|
# Preserve passthru so mkVirtualEnv can resolve dependency groups.
|
|
# Copy .pyi stub + py.typed marker so basedpyright can find the types.
|
|
exo-pyo3-bindings = pkgs.stdenv.mkDerivation {
|
|
pname = "exo-pyo3-bindings";
|
|
version = "0.1.0";
|
|
src = self'.packages.exo_pyo3_bindings;
|
|
# Install from pre-built wheel
|
|
nativeBuildInputs = [ final.pyprojectWheelHook ];
|
|
dontStrip = true;
|
|
passthru = prev.exo-pyo3-bindings.passthru or { };
|
|
postInstall = ''
|
|
local siteDir=$out/${final.python.sitePackages}/exo_pyo3_bindings
|
|
cp ${inputs.self}/rust/exo_pyo3_bindings/exo_pyo3_bindings.pyi $siteDir/
|
|
touch $siteDir/py.typed
|
|
'';
|
|
};
|
|
};
|
|
buildSystemsOverlay = final: prev: { } //
|
|
lib.optionalAttrs isDarwin
|
|
{
|
|
mlx = prev.mlx.overrideAttrs (old:
|
|
let
|
|
# Static dependencies included directly during compilation
|
|
gguf-tools = pkgs.fetchFromGitHub {
|
|
owner = "antirez";
|
|
repo = "gguf-tools";
|
|
rev = "8fa6eb65236618e28fd7710a0fba565f7faa1848";
|
|
hash = "sha256-15FvyPOFqTOr5vdWQoPnZz+mYH919++EtghjozDlnSA=";
|
|
};
|
|
|
|
metal_cpp = pkgs.fetchzip {
|
|
url = "https://developer.apple.com/metal/cpp/files/metal-cpp_26.zip";
|
|
hash = "sha256-7n2eI2lw/S+Us6l7YPAATKwcIbRRpaQ8VmES7S8ZjY8=";
|
|
};
|
|
|
|
nanobind = pkgs.fetchFromGitHub {
|
|
owner = "wjakob";
|
|
repo = "nanobind";
|
|
rev = "v2.10.2";
|
|
hash = "sha256-io44YhN+VpfHFWyvvLWSanRgbzA0whK8WlDNRi3hahU=";
|
|
fetchSubmodules = true;
|
|
};
|
|
in
|
|
{
|
|
nativeBuildInputs = (old.nativeBuildInputs or [ ]) ++ [ pkgs.cmake self'.packages.metal-toolchain ];
|
|
# TODO: non-sdk_26 support
|
|
buildInputs = (old.buildInputs or [ ])
|
|
++ [ gguf-tools pkgs.fmt pkgs.nlohmann_json pkgs.apple-sdk_26 ];
|
|
patches = [
|
|
(pkgs.replaceVars ../nix/darwin-build-fixes.patch {
|
|
sdkVersion = pkgs.apple-sdk_26.version;
|
|
inherit (self'.packages.metal-toolchain) metalVersion;
|
|
})
|
|
];
|
|
postPatch = ''
|
|
substituteInPlace mlx/backend/cpu/jit_compiler.cpp \
|
|
--replace-fail "g++" "${lib.getExe' pkgs.stdenv.cc "c++"}"
|
|
'';
|
|
|
|
DEV_RELEASE = 1;
|
|
CMAKE_ARGS = toString ([
|
|
(lib.cmakeBool "USE_SYSTEM_FMT" true)
|
|
(lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_GGUFLIB" "${gguf-tools}")
|
|
(lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_JSON" "${pkgs.nlohmann_json.src}")
|
|
(lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_NANOBIND" "${nanobind}")
|
|
(lib.cmakeBool "FETCHCONTENT_FULLY_DISCONNECTED" true)
|
|
(lib.cmakeBool "MLX_BUILD_CPU" true)
|
|
(lib.cmakeBool "MLX_BUILD_METAL" true)
|
|
(lib.cmakeOptionType "string" "CMAKE_INSTALL_LIBDIR" "lib")
|
|
(lib.cmakeOptionType "filepath" "FETCHCONTENT_SOURCE_DIR_METAL_CPP" "${metal_cpp}")
|
|
(lib.cmakeOptionType "string" "CMAKE_OSX_DEPLOYMENT_TARGET" "${pkgs.apple-sdk_26.version}")
|
|
(lib.cmakeOptionType "filepath" "CMAKE_OSX_SYSROOT" "${pkgs.apple-sdk_26.passthru.sdkroot}")
|
|
] ++ lib.optionals (isDarwin && isx86_64) [
|
|
(lib.cmakeBool "MLX_ENABLE_X64_MAC" true)
|
|
]);
|
|
SDKROOT = pkgs.apple-sdk_26.passthru.sdkroot;
|
|
MACOSX_DEPLOYMENT_TARGET = pkgs.apple-sdk_26.version;
|
|
});
|
|
} // lib.optionalAttrs isLinux {
|
|
mlx = prev.mlx.overrideAttrs (old: {
|
|
buildInputs = old.buildInputs ++ lib.optionals cudaSupport cudaLibs;
|
|
autoPatchelfIgnoreMissingDeps = lib.optionals cudaSupport [ "libcuda.so.1" ];
|
|
postInstall = (old.postInstall or "") + ''
|
|
cp -r "${final.${libmlx_source}}/${final.python.sitePackages}/mlx" "$out/${final.python.sitePackages}/mlx/"
|
|
'';
|
|
});
|
|
} // lib.optionalAttrs cudaSupport {
|
|
"${libmlx_source}" = prev."${libmlx_source}".overrideAttrs (old: {
|
|
buildInputs = old.buildInputs ++ cudaLibs;
|
|
autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ];
|
|
});
|
|
nvidia-cufile = prev.nvidia-cufile.overrideAttrs (old: {
|
|
buildInputs = old.buildInputs ++ [ pkgs.rdma-core ];
|
|
autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ];
|
|
});
|
|
nvidia-cusolver = prev.nvidia-cusolver.overrideAttrs (old: {
|
|
buildInputs = old.buildInputs ++ cudaLibs;
|
|
autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ];
|
|
});
|
|
nvidia-nvshmem-cu13 = prev.nvidia-nvshmem-cu13.overrideAttrs (old: {
|
|
buildInputs = old.buildInputs ++ [ pkgs.rdma-core pkgs.pmix pkgs.libfabric pkgs.ucx pkgs.openmpi ];
|
|
autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ];
|
|
});
|
|
nvidia-cusparse = prev.nvidia-cusparse.overrideAttrs (old: {
|
|
buildInputs = old.buildInputs ++ [ cudaLibs ];
|
|
autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ];
|
|
});
|
|
torch = prev.torch.overrideAttrs (old: {
|
|
buildInputs = old.buildInputs ++ cudaLibs;
|
|
autoPatchelfIgnoreMissingDeps = [ "libcuda.so.1" ];
|
|
});
|
|
};
|
|
pyprojectOverlay = workspace.mkPyprojectOverlay {
|
|
sourcePreference = "wheel";
|
|
dependencies = { exo = [ uv_extra ]; exo-bench = [ ]; };
|
|
};
|
|
editableOverlay = workspace.mkEditablePyprojectOverlay {
|
|
# Use environment variable pointing to editable root directory
|
|
root = "$REPO_ROOT";
|
|
members = [ "exo" "exo-bench" ];
|
|
};
|
|
pythonSet = (pkgs.callPackage inputs.pyproject-nix.build.packages {
|
|
inherit python;
|
|
}).overrideScope (
|
|
lib.composeManyExtensions [
|
|
inputs.pyproject-build-systems.overlays.default
|
|
pyprojectOverlay
|
|
exoOverlay
|
|
buildSystemsOverlay
|
|
]
|
|
);
|
|
|
|
mkApp = cmd: name: members: pkgs.writeShellApplication {
|
|
inherit name;
|
|
runtimeEnv = {
|
|
EXO_DASHBOARD_DIR = self'.packages.dashboard;
|
|
EXO_RESOURCES_DIR = inputs.self + /resources;
|
|
};
|
|
runtimeInputs = [
|
|
# mlx and mlx-cuda ship clashing cmake files - we dont need them at runtime anyway
|
|
((pythonSet.mkVirtualEnv "${name}-env" members).overrideAttrs (_: { venvSkip = [ "lib/python${python.pythonVersion}/site-packages/mlx/share/cmake/*" ]; }))
|
|
]
|
|
++ lib.optionals isDarwin [ pkgs.macmon ];
|
|
text = "exec " + lib.optionalString cudaSupport "${lib.getExe pkgs.nix-gl-host} " + cmd;
|
|
};
|
|
in
|
|
{
|
|
inherit pythonSet;
|
|
editablePythonSet = pythonSet.overrideScope editableOverlay;
|
|
mkPythonScript = members: name: path: mkApp ''python ${path} "$@"'' name members;
|
|
mkExo = name: members: mkApp ''exo "$@"'' name members;
|
|
};
|
|
in
|
|
{
|
|
perSystem =
|
|
{ self', pkgs, unfreePkgs, lib, ... }:
|
|
let
|
|
inherit (pkgs.stdenv.hostPlatform) isLinux;
|
|
inherit (mkPythonSet { inherit self' pkgs lib; }) pythonSet editablePythonSet mkPythonScript mkExo;
|
|
|
|
exoVenv = pythonSet.mkVirtualEnv "exo-env" { exo = lib.optionals isLinux [ "cpu" ]; };
|
|
|
|
# Virtual environment with dev dependencies for testing
|
|
testVenv = pythonSet.mkVirtualEnv "exo-test-env" {
|
|
exo = [ "dev" ] ++ lib.optionals isLinux [ "cpu" ]; # Include pytest, pytest-asyncio, pytest-env
|
|
};
|
|
|
|
mkBenchScript = mkPythonScript { exo-bench = [ ]; };
|
|
|
|
mkSimplePythonScript = name: path: pkgs.writeShellApplication {
|
|
inherit name;
|
|
runtimeInputs = [ pkgs.python313 ];
|
|
text = ''exec python ${path} "$@"'';
|
|
};
|
|
|
|
in
|
|
{
|
|
packages = {
|
|
exo = mkExo "exo" { exo = lib.optionals isLinux [ "cpu" ]; };
|
|
# for devShell
|
|
exo-venv = exoVenv;
|
|
editableVenv = editablePythonSet.mkVirtualEnv "exo-dev-env" { exo = [ "dev" ]; };
|
|
# for running tests in ci
|
|
exo-test-env = testVenv;
|
|
exo-bench = mkBenchScript "exo-bench" (inputs.self + /bench/exo_bench.py);
|
|
exo-eval = mkBenchScript "exo-eval" (inputs.self + /bench/exo_eval.py);
|
|
exo-eval-tool-calls = mkBenchScript "exo-eval-tool-calls" (inputs.self + /bench/eval_tool_calls.py);
|
|
# used by ./tests/run_exo_on.sh
|
|
exo-get-all-models-on-cluster = mkSimplePythonScript "exo-get-all-models-on-cluster" (inputs.self + /tests/get_all_models_on_cluster.py);
|
|
} // lib.optionalAttrs isLinux {
|
|
exo-cuda-12 = (mkPythonSet { inherit self' lib; inherit (unfreePkgs.pkgsCuda.cudaPackages_12) pkgs; }).mkExo "exo-cuda-12" { exo = [ "cuda12" ]; };
|
|
exo-cuda-13 = (mkPythonSet { inherit self' lib; inherit (unfreePkgs.pkgsCuda.cudaPackages_13) pkgs; }).mkExo "exo-cuda-13" { exo = [ "cuda13" ]; };
|
|
};
|
|
|
|
checks = {
|
|
lint = pkgs.runCommand "ruff-lint" { } ''
|
|
export RUFF_CACHE_DIR="$TMPDIR/ruff-cache"
|
|
${pkgs.ruff}/bin/ruff check ${inputs.self}
|
|
touch $out
|
|
'';
|
|
|
|
typecheck = pkgs.runCommand "typecheck"
|
|
{
|
|
nativeBuildInputs = [
|
|
testVenv
|
|
pkgs.basedpyright
|
|
];
|
|
}
|
|
''
|
|
cd ${inputs.self}
|
|
export HOME=$TMPDIR
|
|
basedpyright --pythonpath ${testVenv}/bin/python --project ${inputs.self}/pyproject.toml
|
|
touch $out
|
|
'';
|
|
};
|
|
};
|
|
}
|