Skip to main content

boatramp_node/
node.rs

1//! The node-graph assembly: given a built store (blobs + KV), a configured
2//! [`Auth`](boatramp_server::Auth), and resolved
3//! [`ServerOptions`](boatramp_server::ServerOptions), wire the deploy store,
4//! handler runtime, compute reconcile loop, and domain-verify reconcile loop
5//! into a [`RunningNode`] ready to hand to a transport (`serve_with` & friends).
6//!
7//! This is the headline extraction of `PLAN-node-library`: the binary's
8//! `serve::run` used to inline this wiring, so no embedder or in-process test
9//! could exercise the same graph the `boatramp serve` binary runs. `run` now
10//! resolves the *environment* (args -> backends -> store, signal handlers,
11//! migration, auth) and calls [`assemble`]; the cluster path keeps its own inline
12//! copy until a later step converges it here.
13
14use std::path::Path;
15use std::sync::Arc;
16
17use boatramp_core::deploy::DeployStore;
18use boatramp_core::kv::KvStore;
19use boatramp_core::Storage;
20
21use crate::config::ServerConfig;
22use crate::error::{Error, Result};
23
24/// How often the compute reconcile loop converges desired vs actual workloads.
25/// Defaults to 30s; override with `BOATRAMP_COMPUTE_RECONCILE_TICK_MS` (milliseconds)
26/// so compute-backed tests can converge in a fraction of a second instead of
27/// waiting a full tick for the launch/scale reconcile.
28pub fn compute_reconcile_tick() -> std::time::Duration {
29    std::env::var("BOATRAMP_COMPUTE_RECONCILE_TICK_MS")
30        .ok()
31        .and_then(|s| s.parse::<u64>().ok())
32        .filter(|&ms| ms > 0)
33        .map(std::time::Duration::from_millis)
34        .unwrap_or(std::time::Duration::from_secs(30))
35}
36/// How often the domain-verify reconcile loop re-checks pending challenges.
37pub const DOMAIN_VERIFY_RECONCILE_TICK: std::time::Duration = std::time::Duration::from_secs(60);
38/// How long a compute workload may be idle before scale-to-zero sleeps it.
39pub const COMPUTE_IDLE_TIMEOUT: std::time::Duration = std::time::Duration::from_secs(300);
40
41/// The built store + resolved config handed to [`assemble`]. Owns the blob/KV
42/// backends and the auth/options the caller already resolved; borrows the parsed
43/// config and data directory.
44pub struct NodeInput<'a> {
45    /// The full parsed server config (the handler + compute sections are read here).
46    pub config: &'a ServerConfig,
47    /// The node data directory (per-site SQL, handler state).
48    pub data_dir: &'a Path,
49    /// The object store built by [`crate::blobs::build_blobs`].
50    pub storage: Arc<dyn Storage>,
51    /// The metadata KV, already cache-fronted, built by [`crate::backends::build_kv`].
52    pub kv: Arc<dyn KvStore>,
53    /// The control-plane auth built by [`crate::auth::configure_auth`].
54    pub auth: boatramp_server::Auth,
55    /// Server options, already carrying the resolved posture, daemon runtime, and
56    /// (post-`configure_auth`/`configure_oidc`) issuer / OIDC verifier.
57    pub options: boatramp_server::ServerOptions,
58    /// The public HTTP serve bind address, if known — used (under
59    /// `allow_guest_self_egress`) to let a handler guest's `wasi:http` reach this
60    /// instance's own front door over loopback. `None` (an in-process embedder with no
61    /// listener) disables self-egress.
62    pub serve_addr: Option<std::net::SocketAddr>,
63    /// The cloud blob-change watch provider (FA-5b2), if the backend is a cloud one.
64    pub watch_provider: Option<Arc<dyn boatramp_core::blob_provision::WatchProvider>>,
65    /// The provisioning tier for the watch provider.
66    pub provision_tier: boatramp_core::blob_notify::ProvisionTier,
67    /// The `wasi:messaging` substrate override for the handler runtime. `None` uses
68    /// the single-node default (`LogMessaging` over the same backends); the cluster
69    /// path passes its Raft-backed coordinator.
70    pub messaging: Option<Arc<dyn boatramp_core::messaging::Messaging>>,
71    /// The single leader gate for cron firing + the compute / domain-verify reconcile
72    /// loops. Single-node passes an always-true gate (there is one node); the cluster
73    /// passes its Raft `is_leader` check so a single node drives each sweep.
74    pub is_leader: boatramp_server::CronLeaderGate,
75    /// This node's compute scheduler id (`0` single-node; the cluster node id in a
76    /// fleet, so replicas are tagged to the right node).
77    pub node_id: u64,
78    /// The binary the re-exec'd compute workers run as — the container backend's
79    /// `__sandbox` jailer and the microVM backends' `__vmm-run`/`__vz-run` VM hosts.
80    /// `None` uses this process's own executable (`current_exe`), which is what
81    /// `boatramp serve` wants (the child *is* boatramp). An **embedding harness**
82    /// whose own binary doesn't implement those subcommands should point this at a
83    /// built `boatramp` binary, so it can drive the real container/microVM backends
84    /// in-process (only the per-workload worker re-execs; the serving plane stays
85    /// embedded). The docker backend needs neither — it talks to a daemon.
86    pub worker_exe: Option<std::path::PathBuf>,
87}
88
89/// A fully wired node: the deploy store, handler runtime, auth, and options a
90/// transport consumes, plus the detached reconcile loops kept alive for the
91/// node's serving life. Destructure it and hold `reconcile` across the serve
92/// await so the loops outlive assembly.
93pub struct RunningNode {
94    /// The deploy store (blob + KV) the router serves from.
95    pub deploy: DeployStore,
96    /// The handler runtime for wasm handlers (a disabled build ⇒ a no-op runtime).
97    pub handlers: boatramp_server::HandlerRuntime,
98    /// The control-plane auth.
99    pub auth: boatramp_server::Auth,
100    /// The resolved server options.
101    pub options: boatramp_server::ServerOptions,
102    /// The detached reconcile loops (compute + domain-verify). Tokio `JoinHandle`s
103    /// do not abort on drop, so the loops run for the process life regardless; the
104    /// handles are retained so an embedder can join/abort them on shutdown.
105    pub reconcile: Vec<tokio::task::JoinHandle<()>>,
106}
107
108/// The instance's own serve socket(s) a guest self-call may reach, given the bind `addr` and
109/// whether the posture (`allow_guest_self_egress`) permits it. A wildcard bind
110/// (`0.0.0.0`/`::`) is reachable over loopback, so it normalizes to `127.0.0.1` **and** `::1`
111/// on the serve port; a specific bind is reachable at itself. Empty when disabled or no
112/// listener.
113fn self_egress_addrs(
114    addr: Option<std::net::SocketAddr>,
115    enabled: bool,
116) -> Vec<std::net::SocketAddr> {
117    use std::net::{IpAddr, Ipv4Addr, Ipv6Addr, SocketAddr};
118    let Some(addr) = addr.filter(|_| enabled) else {
119        return Vec::new();
120    };
121    if addr.ip().is_unspecified() {
122        let port = addr.port();
123        vec![
124            SocketAddr::new(IpAddr::V4(Ipv4Addr::LOCALHOST), port),
125            SocketAddr::new(IpAddr::V6(Ipv6Addr::LOCALHOST), port),
126        ]
127    } else {
128        vec![addr]
129    }
130}
131
132/// Whether a node with **no configured control-plane issuer** should auto-provision an ephemeral
133/// in-memory fleet signer (for host session cookies + delegable capabilities) from its serve bind.
134///
135/// The fleet signer is a distinct trust domain from control-plane admin auth, so a DEV / loopback
136/// node with auth disabled should still be able to sign/verify session cookies and capabilities. It
137/// is gated **strictly to a loopback bind** (`127.0.0.1`/`::1`) or an in-process embedder with no
138/// bind address (`None`): NEVER a public or wildcard (`0.0.0.0`/`::`, reachable off-host) bind,
139/// where an ephemeral key would silently invalidate live capabilities across a restart — a
140/// public node that wants guest capabilities without control-plane auth must supply a persistent
141/// key. (`is_loopback()` is already `false` for a wildcard/unspecified address, so a `0.0.0.0` bind
142/// correctly does NOT auto-provision.)
143#[cfg(feature = "handlers")]
144fn should_autoprovision_fleet_signer(serve_addr: Option<std::net::SocketAddr>) -> bool {
145    serve_addr.is_none_or(|a| a.ip().is_loopback())
146}
147
148/// Wire [`NodeInput`] into a [`RunningNode`]: build the handler runtime, the
149/// deploy store (materializing the reserved `default` project), the compute
150/// backends + reconcile loop, and the domain-verify reconcile loop.
151///
152/// The caller has already built the store and configured auth/OIDC on `options`;
153/// this is the pure node-graph wiring, identical to what `boatramp serve` runs.
154pub async fn assemble(input: NodeInput<'_>) -> Result<RunningNode> {
155    let NodeInput {
156        config,
157        data_dir,
158        storage,
159        kv,
160        auth,
161        options,
162        serve_addr,
163        watch_provider,
164        provision_tier,
165        messaging,
166        is_leader,
167        node_id,
168        worker_exe,
169    } = input;
170    // Copy out the posture scalars up front so `options` can be moved into the
171    // returned `RunningNode` without a lingering borrow.
172    let max_handler_blob_bytes = options.posture.max_handler_blob_bytes;
173    let max_component_bytes = options.posture.max_component_bytes;
174    let allow_guest_private_egress = options.posture.allow_guest_private_egress;
175    let allow_env_secret_refs = options.posture.allow_env_secret_refs;
176    let allow_guest_email = options.posture.allow_guest_email;
177    // The instance's own serve socket(s) a guest self-call may reach, when the posture allows
178    // it: a wildcard bind (`0.0.0.0`/`::`) is reachable on loopback, so normalize to
179    // `127.0.0.1`/`::1`; a specific bind is itself.
180    let self_egress_addrs = self_egress_addrs(serve_addr, options.posture.allow_guest_self_egress);
181    let allow_shared_kernel = options.posture.allow_shared_kernel_compute;
182    let domain_verify_allow_private = options.posture.domain_verify_allow_private;
183
184    // The deploy store the router serves from — built up front so the handler
185    // runtime's managed compute-backed `sql` binding can resolve DB endpoints from
186    // the same store the reconcile writes.
187    let compute_storage = storage.clone();
188    let deploy = DeployStore::new(storage, kv.clone());
189    // The `[secrets]` envelope (local KEK / Vault) that seals a managed SQL
190    // credential at rest. `None` ⇒ no wrapping (a managed DB then fails closed).
191    let secrets_envelope = build_secrets_envelope(config.secrets.as_ref(), data_dir)?;
192    // The project-scoped internal secret store, built from the same KV + `[secrets]`
193    // envelope that seal managed-DB credentials. Backs both the `boatramp:<name>`
194    // resolver (wired into the handler runtime below, when that feature is present)
195    // and the admin secrets API (threaded into `ServerOptions` unconditionally, so it
196    // works on a lean node too). `None` when no envelope is configured — the admin
197    // endpoints then fail closed with a clear 501, never a panic.
198    let secret_store = secrets_envelope.clone().map(|envelope| {
199        Arc::new(boatramp_core::secret_store::SecretStore::new(
200            kv.clone(),
201            envelope,
202        ))
203    });
204    // The project-scoped SMTP email-profile store, built from the same KV + envelope
205    // (the password is sealed at rest). Backs the admin API (`options` below,
206    // unconditionally, so it works on a lean node) and — when the `email` feature +
207    // `allow_guest_email` posture permit — the runtime's host-side profile
208    // resolution (wired inside `build_handler_runtime`). `None` with no envelope, so
209    // the admin email endpoints fail closed with a clear 501.
210    let email_profile_store = secrets_envelope.clone().map(|envelope| {
211        Arc::new(boatramp_core::email_config::EmailProfileStore::new(
212            kv.clone(),
213            envelope,
214        ))
215    });
216
217    // Dev-posture guest-egress extra CA(s): when the posture permits (off/refused under
218    // multi-tenant) AND the operator pointed `BOATRAMP_GUEST_EGRESS_EXTRA_CA_FILE` at a PEM, parse
219    // it into trust anchors the guest's outbound `wasi:http` TLS client trusts on TOP of the webpki
220    // roots (for a hermetic HTTPS test double). Empty otherwise. A configured-but-unreadable/
221    // unparsable file is a hard config error (fail closed), never a silent no-trust.
222    let guest_egress_extra_roots =
223        load_guest_egress_extra_roots(options.posture.allow_guest_egress_extra_ca)?;
224
225    // The handler runtime reuses the same blob/KV backends (per-site prefixed)
226    // for its wasi:blobstore/keyvalue bindings; the sql binding is selected by
227    // `[handlers.bindings.sql]` (default: per-site libsql files under <data-dir>).
228    let handlers = crate::handlers::build_handler_runtime(
229        kv.clone(),
230        compute_storage.clone(),
231        data_dir,
232        config.handlers.as_ref(),
233        messaging,
234        max_handler_blob_bytes,
235        max_component_bytes,
236        allow_guest_private_egress,
237        self_egress_addrs,
238        guest_egress_extra_roots,
239        allow_env_secret_refs,
240        allow_guest_email,
241        options.posture.require_tenancy_declaration,
242        options.posture.allow_cross_tenant_db,
243        &deploy,
244        secrets_envelope.clone(),
245    )
246    .await?;
247    // Wire the fleet session-cookie signer (R3, PLAN-tenancy-principal): the same issuer that mints
248    // control-plane tokens signs + verifies the host-issued anonymous session cookie AND the
249    // delegable capabilities (PLAN-delegable-capabilities). Handlers-gated: the session-cookie
250    // machinery lives on the handler runtime, so a lean (no-handlers) build has nothing to wire.
251    //
252    // The fleet signer is a DIFFERENT trust domain from control-plane admin auth (signing a
253    // customer's session cookie / an embed capability is not the authority to admit an operator to
254    // the control plane), but production derives it from the control-plane issuer for convenience.
255    // For a DEV / loopback node with control-plane auth disabled (`options.issuer` is `None`),
256    // auto-provision an EPHEMERAL in-memory Ed25519 fleet key so the guest-facing signer just works
257    // — session cookies + capability mint/verify — WITHOUT turning on control-plane auth. Strictly
258    // gated to a loopback bind (or an in-process embedder with no bind address): never on a public
259    // bind, where an ephemeral key would silently invalidate live capabilities across a restart (a
260    // public node that wants guest capabilities without control-plane auth must supply a persistent
261    // key). Ephemeral = issue + verify within one process run; nothing persisted, no cross-process
262    // or cross-deploy trust. Production is byte-identical: a real deploy supplies a control-plane key
263    // ⇒ `issuer` is `Some` ⇒ this fallback is never taken.
264    #[cfg(feature = "handlers")]
265    {
266        let fleet_signer = options.issuer.clone().or_else(|| {
267            should_autoprovision_fleet_signer(serve_addr).then(|| {
268                tracing::warn!(
269                    "control-plane auth is disabled and no signer is configured; auto-provisioning \
270                     an EPHEMERAL in-memory fleet signer (Ed25519) for host session cookies + \
271                     delegable capabilities on this loopback/dev node — regenerated each start, \
272                     never persisted. Configure a control-plane key (or a dedicated signer) for \
273                     production."
274                );
275                Arc::new(boatramp_core::cose::LocalSigner::generate(
276                    boatramp_core::cose::TokenAlg::Ed25519,
277                )) as Arc<dyn boatramp_core::cose::Signer>
278            })
279        });
280        if let Some(issuer) = fleet_signer {
281            handlers.set_session_signer(issuer);
282        }
283    }
284    // Enable guest capability minting (`boatramp:handlers/capability`, PLAN-delegable-capabilities)
285    // when the operator posture allows it. A minted capability is verified against the same fleet
286    // signer as the session cookie (wired just above), so this only enables the mint path + the TTL
287    // ceiling; posture-off (or a zero ceiling) ⇒ not offered (a guest `mint` is access-denied).
288    #[cfg(feature = "capability")]
289    if options.posture.allow_guest_mint_capability {
290        handlers.set_capability_minting(options.posture.max_guest_capability_ttl_secs);
291    }
292    // Per-project tenancy/capability posture overrides (Gap 4a): resolve each
293    // `[security.projects.<p>]` override against the fleet base so one serve process can run a
294    // strict-isolation project beside a looser one on a shared, multi-project machine. Empty ⇒
295    // every project uses the node base wired just above. Only these four in-project knobs are
296    // per-project; cross-project isolation stays structural (project = database).
297    #[cfg(feature = "handlers")]
298    {
299        let base = &options.posture;
300        let overrides: std::collections::BTreeMap<
301            String,
302            boatramp_core::security::ResolvedProjectTenancy,
303        > = config
304            .security
305            .as_ref()
306            .map(|s| {
307                s.projects
308                    .iter()
309                    .map(|(project, ovr)| (project.clone(), base.project_tenancy(ovr)))
310                    .collect()
311            })
312            .unwrap_or_default();
313        handlers.set_project_tenancy_overrides(overrides);
314    }
315    // Wire the guest project self-config capability (`boatramp:handlers/admin`) when the
316    // operator posture enables at least one surface. The controller reuses the same in-process
317    // domain-verify / email-profile / secret / site-config subsystems + the real domain probe;
318    // it's project-scoped per grant and rate-limited + audited. Posture-off ⇒ not offered.
319    #[cfg(feature = "admin")]
320    {
321        use boatramp_handlers::AdminSurface;
322        let p = &options.posture;
323        let mut surfaces = std::collections::BTreeSet::new();
324        if p.allow_guest_admin_domains {
325            surfaces.insert(AdminSurface::Domains);
326        }
327        if p.allow_guest_admin_email {
328            surfaces.insert(AdminSurface::Email);
329        }
330        if p.allow_guest_admin_site {
331            surfaces.insert(AdminSurface::Site);
332        }
333        if p.allow_guest_admin_secrets {
334            surfaces.insert(AdminSurface::Secrets);
335        }
336        if !surfaces.is_empty() {
337            let controller = Arc::new(boatramp_server::ServerAdminController::with_server_probe(
338                deploy.clone(),
339                email_profile_store.clone(),
340                secret_store.clone(),
341                p.domain_verify_allow_private,
342            ));
343            handlers.set_admin(controller, surfaces);
344        }
345    }
346    // Leader-gate cron firing (cluster: only the Raft leader fires; single-node: an
347    // always-true gate, equivalent to the unset default). The same gate drives the
348    // reconcile loops below, so all three converge on one leader per fleet. Only the
349    // handler runtime has a scheduler, so this is a no-op without the `handlers` feature.
350    #[cfg(feature = "handlers")]
351    handlers.set_cron_leader_gate(is_leader.clone());
352    // FA-5b2: on a cloud backend, wire the blob-change notification provisioner +
353    // its tier so adding a `blob` trigger provisions (and removing it retracts).
354    #[cfg(feature = "handlers")]
355    if let Some(provider) = watch_provider {
356        handlers.set_watch_provider(provider);
357        handlers.set_provision_tier(provision_tier);
358    }
359    #[cfg(not(feature = "handlers"))]
360    let _ = (watch_provider, provision_tier);
361
362    // Materialize the reserved `default` project so `project ls` / `project show
363    // default` reflect it on a fresh install, not only after a migration. Best
364    // effort: the reader backstop keeps listings correct even if this write can't
365    // land, so a transient failure must never block serving.
366    match deploy.ensure_default_project().await {
367        Ok(true) => tracing::info!("materialized the reserved `default` project record"),
368        Ok(false) => {}
369        Err(e) => tracing::warn!(
370            error = %e,
371            "could not materialize the `default` project record; readers use the synthesized default"
372        ),
373    }
374    // Wire the function-to-function invoke resolver now the deploy store exists,
375    // so a function granted `invoke` can call a sibling in-process (FI).
376    #[cfg(feature = "handlers")]
377    handlers.set_invoker(deploy.clone());
378
379    // Compute reconcile loop. Single-node is always the "leader". Backends are
380    // built from the `[compute]` config + capability detection; a no-op when none
381    // are registered. Detached for the server's life.
382    let (compute_backends, compute_node) = crate::compute::build_compute(
383        config.compute.as_ref(),
384        compute_storage,
385        data_dir,
386        node_id,
387        !allow_shared_kernel,
388        options.daemon_runtime.clone(),
389        worker_exe.as_deref(),
390    )
391    .await;
392    // Adopt the IPs of already-running replicas into each backend's fresh-on-boot
393    // IP pool BEFORE the reconcile loop starts allocating. A backend with a per-node
394    // pool (the native container backend) rebuilds it empty each process start; without
395    // this the boot reconcile could re-hand a live address to a different workload —
396    // the container-IP collision — or move a replica's endpoint on relaunch. Feeds
397    // every persisted replica's `(workload, replica, endpoint-ip)`; each backend keeps
398    // only the IPs in its own subnet (a cheap no-op for docker/cloudflare/VMM).
399    crate::compute::adopt_running_replica_ips(&deploy, &compute_backends).await;
400    // Per-project internal DNS (service discovery): start the resolver on the bridge
401    // gateway so a guest resolves peers by name within its project. On by default;
402    // starts only when the container backend + bridge are up (Linux). Detached for
403    // the node's serving life (pushed into `reconcile` below). Started before the
404    // reconcile loop consumes `compute_backends` — it borrows the registry to check
405    // the container backend is present.
406    let internal_dns =
407        crate::compute::spawn_internal_dns(config.compute.as_ref(), &compute_backends, &deploy);
408    // Activate the compute sql-shim (PLAN-compute-bindings): bind its listener +
409    // build the resolver when a sql provider and `compute.sql_shim_url` are both present.
410    #[cfg(feature = "handlers")]
411    let sql_resolver = boatramp_server::sql_shim::spawn_sql_shim(
412        handlers.sql_backends(),
413        config.compute.as_ref().and_then(|c| c.sql_shim_url.clone()),
414    )
415    .await;
416    #[cfg(not(feature = "handlers"))]
417    let sql_resolver: Option<Arc<dyn boatramp_core::compute::ComputeBindingResolver>> = None;
418
419    // Managed compute-backed SQL (PLAN-managed-compute-sql P2-b): if the handler
420    // `sql` config declares any managed database, inject its `POSTGRES_*`/`MYSQL_*`
421    // server-init env into the DB workload at launch from the sealed credential.
422    // Reaching here with a managed DB implies an envelope (build_handler_runtime
423    // fails closed otherwise), so the credential store always has one to seal with.
424    // Keep a clone of the secrets envelope for the operator-SQL capability below
425    // (the managed_db_resolver match moves the original).
426    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
427    let operator_envelope = secrets_envelope.clone();
428    // …and a second clone for the tenant-deprovision capability (drops a deleted
429    // tenant's managed DB/role/credential on project/site delete). It needs a real
430    // envelope to seal/unseal + delete per-tenant credentials, so it is wired only
431    // when one is present (same fail-closed gating as the managed-DB paths).
432    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
433    let deprovision_envelope = secrets_envelope.clone();
434    // …and a clone for the provisioning drift-repair capability (owner-model retrofit /
435    // reconcile). Like the migrate path it connects as the sealed owner role and re-seals
436    // credentials, so it needs a real envelope; wired only when one is present.
437    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
438    let repair_envelope = secrets_envelope.clone();
439    // …and a third clone for the soft-delete tombstone reaper (the leader-gated task
440    // that hard-drops a Shared-Postgres tenant once its grace window elapses). It, too,
441    // needs a real envelope to unseal the superuser credential + delete the per-tenant
442    // one on hard-drop.
443    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
444    let reaper_envelope = secrets_envelope.clone();
445    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
446    let managed_db_resolver: Option<Arc<dyn boatramp_core::compute::ManagedDbEnvResolver>> = match (
447        config
448            .handlers
449            .as_ref()
450            .and_then(|h| h.bindings.sql.as_ref()),
451        secrets_envelope,
452    ) {
453        (Some(sql), Some(envelope)) if !sql.databases.is_empty() => {
454            let creds = crate::managed_sql::ManagedSqlCredentials::new(kv.clone(), envelope);
455            let privilege = config
456                .compute
457                .as_ref()
458                .map(|c| c.managed_db_privilege)
459                .unwrap_or_default();
460            let env =
461                crate::managed_sql::ManagedDbEnv::from_config(&sql.databases, creds, privilege);
462            (!env.is_empty()).then(|| Arc::new(env) as Arc<_>)
463        }
464        _ => None,
465    };
466    #[cfg(not(any(feature = "sql-postgres", feature = "sql-mysql")))]
467    let managed_db_resolver: Option<Arc<dyn boatramp_core::compute::ManagedDbEnvResolver>> = None;
468
469    // Turnkey managed DB: auto-register the compute workload(s) backing each managed
470    // co-located database that has none yet, so declaring the `databases` binding is
471    // enough to boot the DB (no separate `compute set` / apply). Tenant-aware — a
472    // `Shared` binding registers its one shared server; a `Single` binding registers
473    // nothing at boot (its per-tenant `<compute>-<ident>` is created durably by the lazy
474    // resolve on first `sql` use and relaunched by the reconcile, so a project that never
475    // uses `sql` — e.g. a static-only site — never gets a spurious DB). Non-clobbering +
476    // idempotent; runs before the reconcile loop so its first tick can launch what it
477    // registered.
478    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
479    if let Some(sql) = config
480        .handlers
481        .as_ref()
482        .and_then(|h| h.bindings.sql.as_ref())
483        .filter(|sql| !sql.databases.is_empty())
484    {
485        crate::managed_sql::auto_register_managed_db_workloads(&deploy, &sql.databases).await;
486    }
487
488    // Operator SQL capability (managed-DB migrations/queries via the sealed credential, resolved
489    // server-side) — backs `POST /api/sql/{db}/{exec,query}`. The SAME concrete NodeOperatorSql
490    // also backs the owner-gated schema-migration runner (which reuses its owner + superuser
491    // backends), so build it ONCE and share it.
492    // Build the SAME concrete NodeOperatorSql once (when a managed DB is configured) and share it:
493    // it backs both `operator_sql` (the sql exec/query cap) and the migration runner (which reuses
494    // its owner + superuser backends). Two separate bindings so neither annotation is a complex type.
495    // The migration substrate is a single dispatcher (crate::managed_sql::DispatchMigrationRunner)
496    // routing each `(project, db)` to its engine's substrate: the sqlx NodeMigrationRunner for a
497    // Postgres/MySQL binding, the LibsqlMigrationRunner for a `libsql` binding. It compiles whenever a
498    // sqlx engine OR `migrate` (⇒ libsql) is on, so the embedded-libsql default can migrate even on a
499    // node with no external sqlx engine. `operator_sql` (the `POST /api/sql/{db}/{exec,query}` cap)
500    // stays sqlx-only — a libsql file has no operator-SQL/credential seam.
501    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
502    let operator_sql: Option<Arc<dyn boatramp_core::sql::OperatorSql>>;
503    // Late-init: the match arms assign it (and, under sqlx, `operator_sql` in the same block), and the
504    // arms differ by feature-cfg — so a direct `let … = match {…}` would need cfg'd arm bodies. The
505    // late-init keeps that readable; the value is always assigned before use.
506    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql", feature = "migrate"))]
507    #[allow(clippy::needless_late_init)]
508    let migration_substrate: Option<Arc<dyn boatramp_core::sql::MigrationSubstrate>>;
509    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql", feature = "migrate"))]
510    match config
511        .handlers
512        .as_ref()
513        .and_then(|h| h.bindings.sql.as_ref())
514        .filter(|sql| !sql.databases.is_empty())
515    {
516        Some(sql) => {
517            // The sqlx (Postgres/MySQL) arm — the shared NodeOperatorSql backs both the operator-SQL
518            // cap and the sqlx migration runner. Only built when a sqlx engine is compiled in.
519            #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
520            let node_op = Arc::new(crate::managed_sql::NodeOperatorSql::new(
521                sql.databases.clone(),
522                kv.clone(),
523                operator_envelope,
524                deploy.clone(),
525            ));
526            // The operator's trusted-extension allowlist — the only extensions a migration may
527            // enable (empty ⇒ none). See ExternalSqlConfig::migrate_trusted_extensions.
528            #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
529            let trusted: std::collections::BTreeSet<String> = sql
530                .migrate_trusted_extensions
531                .clone()
532                .unwrap_or_default()
533                .into_iter()
534                .collect();
535            migration_substrate = Some(Arc::new(crate::managed_sql::DispatchMigrationRunner::new(
536                #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
537                node_op.clone(),
538                #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
539                trusted,
540                #[cfg(feature = "migrate")]
541                sql.databases.clone(),
542            ))
543                as Arc<dyn boatramp_core::sql::MigrationSubstrate>);
544            #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
545            {
546                operator_sql = Some(node_op as Arc<dyn boatramp_core::sql::OperatorSql>);
547            }
548        }
549        None => {
550            #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
551            {
552                operator_sql = None;
553            }
554            migration_substrate = None;
555        }
556    }
557    // When only `migrate` (no sqlx) is compiled, the operator-SQL cap does not exist.
558    #[cfg(all(
559        not(any(feature = "sql-postgres", feature = "sql-mysql")),
560        feature = "migrate"
561    ))]
562    let operator_sql: Option<Arc<dyn boatramp_core::sql::OperatorSql>> = None;
563    #[cfg(not(any(feature = "sql-postgres", feature = "sql-mysql", feature = "migrate")))]
564    let migration_substrate: Option<Arc<dyn boatramp_core::sql::MigrationSubstrate>> = None;
565    #[cfg(not(any(feature = "sql-postgres", feature = "sql-mysql", feature = "migrate")))]
566    let operator_sql: Option<Arc<dyn boatramp_core::sql::OperatorSql>> = None;
567
568    // Tenant-deprovision capability (drop a deleted tenant's managed DB/role/sealed
569    // credential on project/site delete). Wired only when a compute-backed managed
570    // database + a secrets envelope are both present — same gating as operator_sql,
571    // plus the envelope requirement (it must seal/unseal per-tenant credentials).
572    // The soft-delete grace window for a Shared-Postgres managed tenant
573    // (`handlers.bindings.sql.deprovision_grace_secs`, env-settable). Default 7 days;
574    // `0` disables the soft path (immediate hard drop). Threaded to the deprovisioner
575    // (which soft-deletes) and implicitly honored by the reaper (which only ever finds
576    // tombstones a >0 grace produced).
577    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
578    let deprovision_grace_secs = config
579        .handlers
580        .as_ref()
581        .and_then(|h| h.bindings.sql.as_ref())
582        .and_then(|sql| sql.deprovision_grace_secs)
583        .unwrap_or(crate::tenant_sql::DEFAULT_DEPROVISION_GRACE_SECS);
584    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
585    let tenant_deprovisioner: Option<Arc<dyn boatramp_core::sql::TenantDeprovisioner>> = config
586        .handlers
587        .as_ref()
588        .and_then(|h| h.bindings.sql.as_ref())
589        .filter(|sql| !sql.databases.is_empty())
590        .zip(deprovision_envelope)
591        .map(|(sql, envelope)| {
592            Arc::new(crate::tenant_sql::NodeTenantDeprovisioner::new(
593                deploy.clone(),
594                kv.clone(),
595                envelope,
596                sql.databases.clone(),
597                deprovision_grace_secs,
598            )) as Arc<_>
599        });
600    #[cfg(not(any(feature = "sql-postgres", feature = "sql-mysql")))]
601    let tenant_deprovisioner: Option<Arc<dyn boatramp_core::sql::TenantDeprovisioner>> = None;
602
603    // Provisioning drift-repair capability (owner-model retrofit / reconcile) — backs the
604    // `Project·Admin`-gated `/api/repair/{db}` + `/dry-run`. Same gating as operator_sql plus
605    // the envelope requirement (it re-seals the owner credential + connects as it). The
606    // envelope is threaded as `Some(_)` so a lean-but-managed node still gets a clear
607    // per-check error rather than a panic if none is configured.
608    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
609    let tenant_repair: Option<Arc<dyn boatramp_core::sql::TenantRepair>> = config
610        .handlers
611        .as_ref()
612        .and_then(|h| h.bindings.sql.as_ref())
613        .filter(|sql| !sql.databases.is_empty())
614        .map(|sql| {
615            Arc::new(crate::repair::NodeTenantRepair::new(
616                sql.databases.clone(),
617                deploy.clone(),
618                kv.clone(),
619                repair_envelope.clone(),
620            )) as Arc<_>
621        });
622    #[cfg(not(any(feature = "sql-postgres", feature = "sql-mysql")))]
623    let tenant_repair: Option<Arc<dyn boatramp_core::sql::TenantRepair>> = None;
624
625    // Operator compute-exec capability (run a command inside a running workload) —
626    // backs `POST /api/compute/{name}/exec`, gated by the `allow_compute_exec`
627    // posture. Clone the backend registry before the reconcile loop consumes it.
628    let compute_exec: Option<Arc<dyn boatramp_core::compute::ComputeExec>> = Some(Arc::new(
629        crate::compute::NodeComputeExec::new(compute_backends.clone(), deploy.clone()),
630    ) as Arc<_>);
631
632    // Operator volume-reclamation capability (list + remove persistent volumes) —
633    // backs `GET /api/compute/volumes` + `DELETE /api/compute/volumes/{name}`.
634    // Same admin-scoped `/api/compute/*` gate; clone the registry before the
635    // reconcile loop consumes the original below.
636    let compute_volumes: Option<Arc<dyn boatramp_core::compute::ComputeVolumes>> = Some(Arc::new(
637        crate::compute::NodeComputeVolumes::new(compute_backends.clone(), deploy.clone()),
638    )
639        as Arc<_>);
640
641    // Operator reconcile-plane control capability (restart a replica) — backs
642    // `POST /api/compute/maintenance/restart` (admin-scoped). Clone the registry
643    // before the reconcile loop consumes the original below.
644    let compute_control: Option<Arc<dyn boatramp_core::compute::ComputeControl>> = Some(Arc::new(
645        crate::compute::NodeComputeControl::new(compute_backends.clone(), deploy.clone()),
646    )
647        as Arc<_>);
648
649    let compute_reconcile = boatramp_server::spawn_compute_reconcile(
650        deploy.clone(),
651        compute_backends,
652        vec![compute_node],
653        boatramp_core::compute::BackendPolicy::from_shared_kernel_allowed(allow_shared_kernel),
654        is_leader.clone(),
655        compute_reconcile_tick(),
656        COMPUTE_IDLE_TIMEOUT,
657        sql_resolver,
658        managed_db_resolver,
659    );
660
661    // Tenant tombstone reaper: leader-gated hard-drop of soft-deleted Shared-Postgres
662    // tenants past their grace window (safe deprovision — see `tenant_sql`). Wired only
663    // when a compute-backed managed database + a secrets envelope are both present
664    // (same gating as the deprovisioner); each tombstone carries its own server +
665    // superuser, so the reaper needs no per-binding config. A `0` grace never writes a
666    // tombstone, so the sweep is simply inert then.
667    #[cfg(any(feature = "sql-postgres", feature = "sql-mysql"))]
668    let tombstone_reaper: Option<tokio::task::JoinHandle<()>> = config
669        .handlers
670        .as_ref()
671        .and_then(|h| h.bindings.sql.as_ref())
672        .filter(|sql| !sql.databases.is_empty())
673        .zip(reaper_envelope)
674        .map(|(_sql, envelope)| {
675            crate::tenant_sql::spawn_tenant_tombstone_reaper(
676                deploy.clone(),
677                kv.clone(),
678                envelope,
679                is_leader.clone(),
680                crate::tenant_sql::TOMBSTONE_REAPER_TICK,
681            )
682        });
683    #[cfg(not(any(feature = "sql-postgres", feature = "sql-mysql")))]
684    let tombstone_reaper: Option<tokio::task::JoinHandle<()>> = None;
685
686    // Domain-verify auto-complete: periodically re-check every site's pending
687    // ownership challenges and attach any that now pass — a published token (e.g.
688    // via `domain add --provider`) converges without a manual `domain verify`.
689    let dv_reconcile = boatramp_server::spawn_domain_verify_reconcile(
690        deploy.clone(),
691        domain_verify_allow_private,
692        is_leader,
693        DOMAIN_VERIFY_RECONCILE_TICK,
694    );
695
696    // Wire the operator capabilities onto the options the router is built from.
697    let mut options = options;
698    options.operator_sql = operator_sql;
699    options.migration_substrate = migration_substrate;
700    options.tenant_repair = tenant_repair;
701    options.tenant_deprovisioner = tenant_deprovisioner;
702    options.compute_exec = compute_exec;
703    options.compute_volumes = compute_volumes;
704    options.compute_control = compute_control;
705    // The internal secret store backs the admin secrets API (set/list/delete). Not
706    // handlers-gated — it must be reachable even on a lean node.
707    options.secret_store = secret_store;
708    // The email-profile store backs the admin API (`/api/email/profiles`); like the
709    // secret store it is not handlers-gated, so it works on a lean node.
710    options.email_profile_store = email_profile_store;
711
712    // The detached reconcile loops: the always-present compute + domain-verify ones,
713    // plus the optional tenant-tombstone reaper (only when a managed DB is configured).
714    let mut reconcile = vec![compute_reconcile, dv_reconcile];
715    if let Some(reaper) = tombstone_reaper {
716        reconcile.push(reaper);
717    }
718    if let Some(dns) = internal_dns {
719        reconcile.push(dns);
720    }
721
722    Ok(RunningNode {
723        deploy,
724        handlers,
725        auth,
726        options,
727        reconcile,
728    })
729}
730
731/// Build the `[secrets]` envelope (secrets-at-rest wrapping) from `boatramp.cfg`'s
732/// `[secrets]` section: `local` (a machine-local AES-256-GCM KEK) or `vault` (Vault
733/// Env var an operator points at a PEM file of extra CA(s) the guest's outbound `wasi:http` TLS
734/// client should trust on top of the webpki roots — honored only under the
735/// `allow_guest_egress_extra_ca` posture (a hermetic HTTPS test double lever).
736const GUEST_EGRESS_EXTRA_CA_ENV: &str = "BOATRAMP_GUEST_EGRESS_EXTRA_CA_FILE";
737
738/// Parse the operator's guest-egress extra-CA PEM ([`GUEST_EGRESS_EXTRA_CA_ENV`]) into rustls trust
739/// anchors, gated by the `allow_guest_egress_extra_ca` posture (`allow`). No env set ⇒ empty (the
740/// default). `allow == false` (e.g. multi-tenant) with a file set ⇒ empty + a warning (the posture
741/// refuses it). Set + readable + ≥1 cert ⇒ those certs. Set-but-unreadable / no valid cert ⇒ a hard
742/// error (fail closed — a configured-but-broken CA must not silently degrade to no-trust).
743fn load_guest_egress_extra_roots(
744    allow: bool,
745) -> Result<Vec<rustls::pki_types::CertificateDer<'static>>> {
746    let path = match std::env::var(GUEST_EGRESS_EXTRA_CA_ENV) {
747        Ok(p) if !p.is_empty() => p,
748        _ => return Ok(Vec::new()),
749    };
750    if !allow {
751        tracing::warn!(
752            env = GUEST_EGRESS_EXTRA_CA_ENV,
753            "ignoring a guest-egress extra CA: the security posture forbids it \
754             (allow_guest_egress_extra_ca is off — e.g. under multi-tenant)"
755        );
756        return Ok(Vec::new());
757    }
758    let pem =
759        std::fs::read(&path).map_err(|e| Error::GuestEgressCa(format!("reading {path:?}: {e}")))?;
760    let certs = rustls_pemfile::certs(&mut std::io::BufReader::new(&pem[..]))
761        .collect::<std::result::Result<Vec<_>, _>>()
762        .map_err(|e| Error::GuestEgressCa(format!("parsing {path:?}: {e}")))?;
763    if certs.is_empty() {
764        return Err(Error::GuestEgressCa(format!(
765            "{path:?} contained no PEM certificate"
766        )));
767    }
768    tracing::info!(
769        env = GUEST_EGRESS_EXTRA_CA_ENV,
770        count = certs.len(),
771        path = %path,
772        "guest egress trusts operator extra CA(s) (dev-posture; webpki roots still apply)"
773    );
774    Ok(certs)
775}
776
777/// Transit). `None`/empty ⇒ no wrapping. The Vault token is read from the
778/// environment (`token_env`), never a file. This seals a managed SQL credential at
779/// rest; a managed database fails closed without it.
780fn build_secrets_envelope(
781    secrets: Option<&crate::config::SecretsConfig>,
782    data_dir: &Path,
783) -> Result<Option<Arc<dyn boatramp_core::envelope::KeyEnvelope>>> {
784    use boatramp_server::envelope::{build_envelope, EnvelopeSpec};
785    let Some(cfg) = secrets else {
786        return Ok(None);
787    };
788    let spec = match cfg.envelope.as_str() {
789        "" => EnvelopeSpec::None,
790        "local" => EnvelopeSpec::Local {
791            kek_file: cfg
792                .kek_file
793                .clone()
794                .unwrap_or_else(|| data_dir.join("secrets/kek")),
795        },
796        "vault" => {
797            let v = cfg.vault.as_ref().ok_or_else(|| {
798                Error::Envelope(
799                    "secrets.envelope = \"vault\" needs a [secrets.vault] section".into(),
800                )
801            })?;
802            let token = std::env::var(&v.token_env).map_err(|_| {
803                Error::Envelope(format!("Vault token env `{}` is not set", v.token_env))
804            })?;
805            EnvelopeSpec::Vault {
806                addr: v.addr.clone(),
807                key: v.key.clone(),
808                token,
809            }
810        }
811        other => {
812            return Err(Error::Envelope(format!(
813                "unknown secrets.envelope {other:?} (want \"local\" or \"vault\")"
814            )))
815        }
816    };
817    build_envelope(spec).map_err(|e| Error::Envelope(e.to_string()))
818}
819
820#[cfg(all(test, feature = "fs"))]
821mod tests {
822    use super::*;
823    use boatramp_core::kv::MemoryKv;
824    use boatramp_core::security::SecurityProfile;
825
826    /// The dev/loopback ephemeral fleet-signer auto-provision is gated STRICTLY to a loopback bind
827    /// (or an in-process embedder with no bind): a public / wildcard / private-network bind must
828    /// NOT silently provision an ephemeral signing key (it would invalidate live capabilities on a
829    /// restart — such a node must supply a persistent key). This is the security-critical boundary.
830    #[cfg(feature = "handlers")]
831    #[test]
832    fn ephemeral_fleet_signer_auto_provisions_only_on_loopback_or_in_process() {
833        use std::net::SocketAddr;
834        let sa = |s: &str| s.parse::<SocketAddr>().unwrap();
835        // In-process (no listener) and loopback → auto-provision the dev fleet signer.
836        assert!(should_autoprovision_fleet_signer(None));
837        assert!(should_autoprovision_fleet_signer(Some(sa(
838            "127.0.0.1:8080"
839        ))));
840        assert!(should_autoprovision_fleet_signer(Some(sa("[::1]:8080"))));
841        // Off-host-reachable binds MUST NOT auto-provision an ephemeral key: a public IP, a
842        // private-network IP, and a wildcard bind (`0.0.0.0`/`::`, reachable off-host).
843        assert!(!should_autoprovision_fleet_signer(Some(sa(
844            "203.0.113.5:8080"
845        ))));
846        assert!(!should_autoprovision_fleet_signer(Some(sa(
847            "10.0.0.4:8080"
848        ))));
849        assert!(!should_autoprovision_fleet_signer(Some(sa("0.0.0.0:8080"))));
850        assert!(!should_autoprovision_fleet_signer(Some(sa("[::]:8080"))));
851    }
852
853    /// The headline in-process fidelity check (PLAN-node-library N2b.3): `assemble`
854    /// over a temp `FsStorage` + `MemoryKv` produces a `RunningNode` whose deploy
855    /// store is live (the reserved `default` project was materialized during
856    /// assembly) and whose router — the exact one `boatramp serve` builds — answers
857    /// `/healthz`. No listener is bound: the request is driven through the router
858    /// via `tower::oneshot`, so the whole assembly runs in-process.
859    #[tokio::test]
860    async fn assemble_produces_a_serving_node_over_a_temp_store() {
861        use axum::body::Body;
862        use axum::http::{Request, StatusCode};
863        use tower::ServiceExt;
864
865        let tmp = tempfile::tempdir().unwrap();
866        let storage: Arc<dyn Storage> = Arc::new(boatramp_storage::FsStorage::new(tmp.path()));
867        let kv: Arc<dyn KvStore> = Arc::new(MemoryKv::new());
868        let config = ServerConfig::default();
869        let options = boatramp_server::ServerOptions {
870            // The strict `multi-tenant` posture, as an unconfigured `serve` resolves.
871            posture: SecurityProfile::MultiTenant.preset(),
872            ..Default::default()
873        };
874
875        let node = assemble(NodeInput {
876            config: &config,
877            data_dir: tmp.path(),
878            storage,
879            kv,
880            auth: boatramp_server::Auth::disabled(),
881            options,
882            serve_addr: None,
883            watch_provider: None,
884            provision_tier: boatramp_core::blob_notify::ProvisionTier::default(),
885            messaging: None,
886            is_leader: Arc::new(|| true),
887            node_id: 0,
888            worker_exe: None,
889        })
890        .await
891        .expect("assemble a node over a temp store");
892
893        // The deploy store is live: `assemble` already materialized the reserved
894        // `default` project, so a second ensure reports "already present" (`false`).
895        assert!(
896            !node
897                .deploy
898                .ensure_default_project()
899                .await
900                .expect("read the default project"),
901            "assemble should have materialized the default project"
902        );
903
904        // The assembled router (the same wiring `serve` binds) answers /healthz.
905        let router =
906            boatramp_server::router_with(node.deploy, node.auth, node.handlers, node.options);
907        let response = router
908            .oneshot(
909                Request::builder()
910                    .uri("/healthz")
911                    .body(Body::empty())
912                    .unwrap(),
913            )
914            .await
915            .expect("route /healthz");
916        assert_eq!(response.status(), StatusCode::OK);
917    }
918}