swactor/examples/pipeline-parallel-inference/Cargo.toml

43 lines
1.2 KiB
TOML
Raw Normal View History

[workspace]
[package]
name = "pipeline-parallel-inference"
version = "0.1.0"
edition = "2024"
publish = false
[dependencies]
swactor = { path = "../..", features = ["transport", "serde", "std"] }
swactor-process = { path = "../../crates/process" }
serde = { version = "1", features = ["derive"] }
serde_json = "1"
reqwest = { version = "0.12", features = ["json", "stream"] }
futures-util = "0.3"
tokio = { version = "1", features = ["full"] }
distribution = { path = "../../crates/distribution", features = ["iroh"] }
# Live runtime dashboard (HTTP overview/actors/topology pages + the
# datastream-fed Fleet tab, served on localhost when PP_DASHBOARD is set).
dashboard = { path = "../../crates/dashboard" }
iroh = "0.98"
urlencoding = "2"
base64 = "0.22"
libc = "0.2"
[[bin]]
name = "pp-worker"
path = "src/bin/pp_worker.rs"
[[bin]]
name = "pp-orchestrator"
path = "src/bin/pp_orchestrator.rs"
[dev-dependencies]
wiremock = "0.6"
tracing-subscriber = { version = "0.3", features = ["env-filter"] }
# These binaries are shipped to rented GPU nodes over the docker image on a
# cold lease, and over scp through vast.ai's throttled SSH proxy for manual
# swaps. Stripping debug symbols (~26MB -> ~18MB) trims both at no runtime cost.
[profile.release]
strip = true