Skip to main content

mcpmesh_node/
config.rs

1//! The `config.toml` model. Every table and key here is real, implemented surface —
2//! docs/config.md is the operator-facing reference for all of it.
3use figment::{
4    Figment,
5    providers::{Format, Toml},
6};
7use serde::Deserialize;
8use std::collections::BTreeMap;
9use std::path::PathBuf;
10
11#[derive(Debug, Default, Deserialize)]
12#[serde(default)]
13pub struct Config {
14    pub identity: IdentityCfg,
15    pub network: NetworkCfg,
16    pub limits: LimitsCfg,
17    /// Roster-mode `[roster]` tunables: the degraded-expiry grace window, the roster URL +
18    /// poll interval, and the freshness bound — one `RosterState` machine consumes them all.
19    pub roster: RosterCfg,
20    /// `[services.<name>]` registry — each entry is a served MCP server plus its allow
21    /// list. Peers do NOT live in config; they live in the daemon's state store, so
22    /// there is no `[peers]` table here.
23    pub services: std::collections::BTreeMap<String, ServiceCfg>,
24}
25
26/// A `[services.<name>]` entry: exactly one backend kind (`run` xor `socket`) plus the
27/// nicknames/groups admitted to it. The xor is validated at access time via
28/// [`ServiceCfg::backend_result`] rather than at parse time, so a malformed entry is a
29/// per-service error, not a whole-config load failure.
30#[derive(Debug, Default, Deserialize)]
31#[serde(default)]
32pub struct ServiceCfg {
33    /// `run`: spawn this command per session (a stdio MCP server).
34    pub run: Option<Vec<String>>,
35    /// `socket`: dial this local UDS (an already-running MCP server).
36    pub socket: Option<String>,
37    /// STABLE principals admitted to this service (b64u:/eid:/roster names, #38 — never display nicknames).
38    pub allow: Vec<String>,
39    /// Per-service env vars for a `run` backend (#51). The `MCPMESH_PEER_*` identity vars win
40    /// over these. Ignored for a `socket` backend. Default empty.
41    pub env: BTreeMap<String, String>,
42    /// Working directory for a `run` backend (#51). Default: inherit the daemon's cwd.
43    pub cwd: Option<String>,
44}
45
46/// The resolved backend kind of a [`ServiceCfg`], borrowing the config as slices (no
47/// clone). `&[String]`/`&str` rather than `&Vec`/`&String` — idiomatic and gives the
48/// daemon's backend builders the most flexible borrow.
49#[derive(Debug)]
50pub enum Backend<'a> {
51    Run(&'a [String]),
52    Socket(&'a str),
53}
54
55impl ServiceCfg {
56    /// Resolve the backend, enforcing exactly-one-of `run`/`socket`. Both or neither is an
57    /// error — surfaced to the operator, never a silent default.
58    #[allow(dead_code)] // consumed by the daemon service wiring
59    pub fn backend_result(&self) -> Result<Backend<'_>, String> {
60        match (&self.run, &self.socket) {
61            (Some(cmd), None) => Ok(Backend::Run(cmd.as_slice())),
62            (None, Some(p)) => Ok(Backend::Socket(p.as_str())),
63            (Some(_), Some(_)) => Err("service has both run and socket".into()),
64            (None, None) => Err("service has neither run nor socket".into()),
65        }
66    }
67}
68
69#[derive(Debug, Default, Deserialize)]
70#[serde(default)]
71pub struct IdentityCfg {
72    pub device_key: Option<PathBuf>, // None → paths::default_device_key_path()
73    /// This device's suggested name for itself, carried in a minted pairing invite.
74    /// `None` → the daemon defaults to a short fingerprint of the endpoint id.
75    /// Additive (`#[serde(default)]` at the struct level).
76    pub nickname: Option<String>,
77    /// Roster mode: the org id this node joined (pinned at install/join).
78    pub org_id: Option<String>,
79    /// Roster mode: the pinned org-root public key, `b64u:`. The single trust anchor
80    /// roster signatures verify against. Pinned on first roster install / `join`.
81    pub org_root_pk: Option<String>,
82    /// Roster mode: this node's stable user_id in the org. Pinned at `join` (proposed)
83    /// and reconciled to the roster's authoritative value once installed.
84    pub user_id: Option<String>,
85    /// Roster mode: path to this person's user key. Minted by `join`; binds this
86    /// person's devices. `None` → paths::default_user_key_path() when needed.
87    pub user_key: Option<PathBuf>,
88}
89
90/// `[network]`. The knobs are exactly what `daemon::net_plan` implements —
91/// no aspirational surface:
92/// - `relay_mode = "default" | "custom" | "disabled"`. `"custom"` requires `relay_urls`
93///   (self-hosted iroh relays); `"disabled"` is the HERMETIC mode — no relay AND no
94///   discovery (localhost/tests).
95/// - `discovery_mode = "default" | "custom"`. `"custom"` requires `discovery_urls` —
96///   self-hosted pkarr relay URLs (e.g. an iroh-dns-server), used for BOTH publishing and
97///   resolving peer addresses in place of n0's DNS/pkarr. Ignored (off) when
98///   `relay_mode = "disabled"`.
99///
100/// Unknown modes or a `custom` without URLs are startup ERRORS (`net_plan`), never a silent
101/// fallback — a metadata-privacy knob must not quietly revert to public infrastructure.
102#[derive(Debug, Clone, Deserialize)]
103#[serde(default)]
104pub struct NetworkCfg {
105    pub relay_mode: String,
106    /// Self-hosted relay URLs, required when `relay_mode = "custom"`.
107    pub relay_urls: Vec<String>,
108    pub discovery_mode: String,
109    /// Self-hosted pkarr relay URLs, required when `discovery_mode = "custom"`.
110    pub discovery_urls: Vec<String>,
111    /// TESTING ONLY (#116): force application data over the RELAY even when a direct path exists.
112    ///
113    /// Requires the `unstable-relay-only` cargo feature. Without it this field still PARSES — a
114    /// config must stay portable between a test build and a production one — but is ignored with a
115    /// `warn!`. It is never a startup error: a testing switch must not brick a node, and it must
116    /// never be ignored SILENTLY, because believing you tested the relay when you did not is the
117    /// exact failure #116 reports.
118    ///
119    /// Selects the relay path; it does NOT prevent hole-punching (that is socket-level behaviour a
120    /// `PathSelector` cannot reach). A direct path may still form — it simply never carries data,
121    /// and `status` reports `relay` because #64 derives the path from `is_selected()`.
122    pub relay_only: bool,
123    /// QUIC idle timeout in seconds (#56) — how long a connection survives with NO traffic and no
124    /// keepalive before the transport closes it. `None` = iroh's default, **30s** on iroh 1.0.3.
125    ///
126    /// This is not "how long an idle session lives". iroh keepalives every 5s by default, so a held
127    /// session survives indefinitely while the process runs; this is what detects a peer that
128    /// VANISHED.
129    ///
130    /// **It is NEGOTIATED, not imposed.** QUIC takes the MINIMUM of the two peers' advertised
131    /// values (RFC 9000 §10.1), so raising this on one node achieves nothing against a peer still
132    /// on the default — the connection still times out at 30s. Raising it is only meaningful when
133    /// every node is configured together; lowering it works one-sidedly.
134    ///
135    /// `0` means "no timeout" from THIS side, which likewise yields the peer's value; against a
136    /// default peer that is still 30s. Only if both sides say `0` does a vanished peer go
137    /// undetected at the transport layer.
138    #[serde(default)]
139    pub idle_timeout_secs: Option<u64>,
140    /// QUIC keepalive interval in seconds (#56) — how often the transport PINGs an otherwise idle
141    /// connection. `None` = iroh's default, **5s** on iroh 1.0.3.
142    ///
143    /// Sets BOTH the connection-level and the per-path keepalive — setting only the former would
144    /// leave every path pinging at iroh's 5s regardless.
145    ///
146    /// A transport keepalive carries no method-bearing frame, so it does NOT consume a
147    /// `[limits].rate_limit_per_min` token — unlike an application-level heartbeat, which does.
148    ///
149    /// Must be less than the EFFECTIVE idle timeout — `idle_timeout_secs` if set, otherwise iroh's
150    /// 30s — or boot fails. Note that effective timeout is the negotiated minimum, so a value that
151    /// passes this check locally can still be too slow for a peer with a shorter one.
152    ///
153    /// **This can only LOWER the ping rate.** iroh caps the per-path keepalive at 5s and silently
154    /// discards anything larger, so a value above 5 would leave every path pinging at 5s anyway —
155    /// boot refuses it rather than pretend it took effect. There is no supported way to reduce
156    /// keepalive traffic on a metered link with iroh 1.0.3.
157    #[serde(default)]
158    pub keep_alive_secs: Option<u64>,
159}
160impl Default for NetworkCfg {
161    fn default() -> Self {
162        Self {
163            relay_mode: "default".into(),
164            relay_urls: Vec::new(),
165            discovery_mode: "default".into(),
166            discovery_urls: Vec::new(),
167            relay_only: false,
168            idle_timeout_secs: None,
169            keep_alive_secs: None,
170        }
171    }
172}
173
174/// `[limits]`. NOTE — the frame cap is deliberately NOT here: the 16 MiB `max_frame`
175/// default is a fixed CONSTANT at each wire (`mcpmesh_net::endpoint` for the mesh,
176/// `ipc::MAX_FRAME_BYTES` for the control socket, `backends::MAX_FRAME_BYTES` for local MCP
177/// servers), not a config tunable. A `max_frame` config field existed historically but was never
178/// threaded into any `FrameReader` (dead surface); threading it into the mesh path would widen
179/// `mcpmesh-net`'s public API for no demonstrated need, so the field was removed instead (serde
180/// ignores an unknown `max_frame` key in existing configs).
181#[derive(Debug, Deserialize)]
182#[serde(default)]
183pub struct LimitsCfg {
184    pub rate_limit_per_min: u32,
185    pub max_inflight: u32,
186    pub max_sessions: u32,
187    /// Per-authenticated-endpoint app-blob BYTE budget, bytes per minute (#84a).
188    ///
189    /// **0 = unlimited, and that is the default**, so an existing deployment is unchanged on
190    /// upgrade. The pre-existing blob limiter counts CONNECTIONS, which cannot see one granted
191    /// peer re-pulling a 4 GB blob on each of 60 connections a minute; this bounds the bytes.
192    ///
193    /// A peer that exceeds it gets its transfer ABORTED (retryable), not paced — pacing holds the
194    /// request open and turns a bandwidth problem into an unbounded-concurrency one.
195    ///
196    /// **Use 0 or at least 32768** (two chunks); a value in `1..32768` is FLOORED to 32768.
197    ///
198    /// Admission reserves one chunk before any bytes and the transfer then meters its own chunks,
199    /// so a sub-floor budget does not fail closed — it silently caps every servable blob at
200    /// roughly `budget - 16384` bytes and truncates anything larger. Measured: 20480 serves a
201    /// 4 KiB blob and nothing bigger. Two earlier drafts of this comment got that wrong, first
202    /// recommending the bricking value and then claiming it failed closed.
203    ///
204    /// Requires a restart: the limiter and the provider's event mask are both built once at boot.
205    pub blob_bytes_per_min: u64,
206    /// Audit-log retention window in calendar months (#88). **0 = keep forever, and that is the
207    /// default** — flipping today's keep-everything behavior to auto-deletion is a product call,
208    /// deliberately not made here. When N > 0, boot deletes monthly audit files older than the
209    /// last N months (the current month counts as month 1). Boot-time only: a long-running
210    /// daemon prunes on its next start; the `audit_prune` verb covers live needs.
211    pub audit_retain_months: u32,
212}
213impl Default for LimitsCfg {
214    fn default() -> Self {
215        Self {
216            rate_limit_per_min: 120,
217            max_inflight: 16,
218            max_sessions: 4,
219            blob_bytes_per_min: 0, // unlimited: opt-in, no behaviour change on upgrade
220            audit_retain_months: 0, // keep forever: opt-in, no behaviour change on upgrade
221        }
222    }
223}
224
225/// The default degraded-expiry grace window (`[roster].grace_period` default "72h").
226/// A stale roster keeps serving for this window past `expires_at` (with a warning) before it
227/// stops granting roster identity. Kept here so [`RosterCfg::default`] and the parse fallback
228/// share one source; the gate mirrors it as `roster::gate::DEFAULT_GRACE_SECS`.
229const DEFAULT_GRACE_SECS: i64 = 72 * 3600;
230
231/// The default freshness bound (`[roster].max_staleness`, default "24h" = 86400s). A roster
232/// this node has not re-confirmed current within this window degrades on the SAME `RosterState`
233/// machine as expiry (warnings within `grace`, then serving stops) — bounding adversarial staleness at
234/// `max_staleness + grace` independent of `expires_at`. Shared by [`RosterCfg::default`] + the parse
235/// fallback.
236const DEFAULT_MAX_STALENESS_SECS: i64 = 24 * 3600;
237
238/// The `[roster]` config table. `grace_period` is the degraded-expiry grace window — how
239/// long a roster past `expires_at` keeps serving (degraded, warning) before it stops. Additive
240/// (`#[serde(default)]`): a config with no `[roster]` table gets the 72h default.
241#[derive(Debug, Deserialize)]
242#[serde(default)]
243pub struct RosterCfg {
244    /// Degraded-expiry grace window: `"72h"` / `"24h"` / plain seconds (default "72h").
245    pub grace_period: String,
246    /// The pinned roster URL for the HTTPS poll. Operator-managed static hosting; also how a
247    /// joiner bootstraps its FIRST roster. `None` → no URL poll (manual installs only).
248    /// Additive (`#[serde(default)]`): a config with no `url` key gets `None`.
249    pub url: Option<String>,
250    /// How often to poll `url` (default "1h"). Total-parse like `grace_period` — an
251    /// unparseable value falls back to the hourly default rather than disabling the poll.
252    pub poll_interval: String,
253    /// The freshness bound (default "24h"): how long this node may go without re-confirming
254    /// the installed roster current (via a TLS URL poll ≥ installed, a gossip install, or a
255    /// manual install) before it degrades on the SAME `RosterState` machine as expiry. Total-parse
256    /// like `grace_period` (an unparseable value falls back to the 24h default — a typo never disables
257    /// the bound). Additive (`#[serde(default)]`): a config with no `max_staleness` key gets 24h.
258    pub max_staleness: String,
259}
260impl Default for RosterCfg {
261    fn default() -> Self {
262        Self {
263            grace_period: "72h".into(),
264            url: None,
265            poll_interval: "1h".into(),
266            max_staleness: "24h".into(),
267        }
268    }
269}
270
271impl RosterCfg {
272    /// The grace window in SECONDS. An absent or unparseable `grace_period` falls back to the 72h
273    /// default rather than erroring — an operator typo must never disable degraded serving, and a
274    /// grace window is advisory, not a security bound (revocation is enforced regardless of
275    /// degraded state).
276    ///
277    /// Two paths degrade on the ONE `RosterState` machine (`RosterView::state`, Approved →
278    /// DegradedGrace → DegradedStopped): expiry (`expires_at` + THIS grace window) and freshness
279    /// (`last_confirmed` + `max_staleness`). Once DegradedStopped, the gate stops granting roster
280    /// identity (fail-closed — revocation is still enforced); within grace, serving continues
281    /// with a warning (`daemon::warn_if_degraded_grace`).
282    pub fn grace_seconds(&self) -> i64 {
283        parse_duration(&self.grace_period).unwrap_or(DEFAULT_GRACE_SECS)
284    }
285
286    /// The URL poll interval in SECONDS (default 3600). Like [`grace_seconds`](Self::grace_seconds)
287    /// it is TOTAL — an absent/unparseable value falls back to the hourly default rather than
288    /// erroring, so an operator typo slows the poll to hourly instead of disabling freshness.
289    pub fn poll_interval_seconds(&self) -> i64 {
290        parse_duration(&self.poll_interval).unwrap_or(3600)
291    }
292
293    /// The freshness bound in SECONDS (default 86400 = 24h). Like [`grace_seconds`](Self::grace_seconds)
294    /// it is TOTAL — an absent/unparseable value falls back to the 24h default rather than erroring, so
295    /// an operator typo tightens/loosens to 24h instead of disabling the freshness bound.
296    pub fn max_staleness_seconds(&self) -> i64 {
297        parse_duration(&self.max_staleness).unwrap_or(DEFAULT_MAX_STALENESS_SECS)
298    }
299}
300
301/// Parse a duration string to SECONDS: a `d`/`h`/`m`/`s` suffix (days/hours/minutes/seconds) or a
302/// bare number (seconds). Trim + suffix-strip + checked multiply; rejects a
303/// negative/overflowing/garbage value as `Err` (the caller supplies the
304/// default). `u64` parse then a checked `i64` conversion: a negative grace is meaningless, so `-1`
305/// fails the `u64` parse and falls back to the default rather than becoming a negative window.
306// Reached only by the accessors above and the `org create --expires` porcelain
307// (`enrollcmd`, the operator-managed validity window — now across the crate seam, hence
308// `pub`; still `#[doc(hidden)]` at the module level). Pure parser — no state.
309pub fn parse_duration(s: &str) -> Result<i64, String> {
310    let s = s.trim();
311    let (num, mult) = if let Some(n) = s.strip_suffix('d') {
312        (n, 24 * 3600)
313    } else if let Some(n) = s.strip_suffix('h') {
314        (n, 3600)
315    } else if let Some(n) = s.strip_suffix('m') {
316        (n, 60)
317    } else if let Some(n) = s.strip_suffix('s') {
318        (n, 1)
319    } else {
320        (s, 1)
321    };
322    num.trim()
323        .parse::<u64>()
324        .ok()
325        .and_then(|v| v.checked_mul(mult))
326        .and_then(|v| i64::try_from(v).ok())
327        .ok_or_else(|| format!("unparseable duration: {s}"))
328}
329
330// figment::Error is ~208 bytes; boxing it would churn the API for a cold path.
331#[allow(clippy::result_large_err)]
332impl Config {
333    #[allow(dead_code)] // exercised by unit tests; config-string entry point for later tooling
334    pub fn from_toml_str(s: &str) -> Result<Self, figment::Error> {
335        Figment::new().merge(Toml::string(s)).extract()
336    }
337
338    /// Missing file → defaults (first run); malformed file → Err.
339    /// Callers must surface the Err — swallowing it silently reverts user choices.
340    pub fn load(path: &std::path::Path) -> Result<Self, figment::Error> {
341        Figment::new().merge(Toml::file(path)).extract()
342    }
343}
344
345#[cfg(test)]
346mod tests {
347    use super::*;
348
349    #[test]
350    fn empty_file_yields_spec_defaults() {
351        let c = Config::from_toml_str("").unwrap();
352        assert_eq!(c.network.relay_mode, "default");
353        assert_eq!(c.network.discovery_mode, "default");
354        assert_eq!(c.limits.rate_limit_per_min, 120);
355        assert_eq!(c.limits.max_inflight, 16);
356        assert_eq!(c.limits.max_sessions, 4);
357    }
358
359    #[test]
360    fn values_override_defaults() {
361        let c = Config::from_toml_str(
362            "[network]\nrelay_mode = \"disabled\"\n[limits]\nrate_limit_per_min = 60\n",
363        )
364        .unwrap();
365        assert_eq!(c.network.relay_mode, "disabled");
366        assert_eq!(c.limits.rate_limit_per_min, 60);
367        assert_eq!(c.limits.max_inflight, 16);
368    }
369
370    /// A legacy config carrying the removed `max_frame` key still loads (serde ignores unknown
371    /// fields) — the frame cap is a fixed constant now, not a tunable (see the `LimitsCfg` doc).
372    #[test]
373    fn legacy_max_frame_key_is_ignored_not_an_error() {
374        let c =
375            Config::from_toml_str("[limits]\nmax_frame = \"1MiB\"\nmax_sessions = 2\n").unwrap();
376        assert_eq!(c.limits.max_sessions, 2);
377    }
378
379    /// The self-hosting knobs parse: `custom` modes with their URL lists. (Validation —
380    /// custom-without-urls, unknown modes — lives in `daemon::net_plan`, tested there.)
381    #[test]
382    fn network_relay_and_discovery_urls_parse() {
383        let c = Config::from_toml_str(
384            "[network]\nrelay_mode = \"custom\"\nrelay_urls = [\"https://relay.acme.com\"]\n\
385             discovery_mode = \"custom\"\ndiscovery_urls = [\"https://dns.acme.com/pkarr\"]\n",
386        )
387        .unwrap();
388        assert_eq!(c.network.relay_mode, "custom");
389        assert_eq!(
390            c.network.relay_urls,
391            vec!["https://relay.acme.com".to_string()]
392        );
393        assert_eq!(c.network.discovery_mode, "custom");
394        assert_eq!(
395            c.network.discovery_urls,
396            vec!["https://dns.acme.com/pkarr".to_string()]
397        );
398        // Absent → empty lists (the defaults need no URLs).
399        let c = Config::from_toml_str("").unwrap();
400        assert!(c.network.relay_urls.is_empty() && c.network.discovery_urls.is_empty());
401    }
402
403    #[test]
404    fn missing_file_loads_defaults() {
405        let dir = tempfile::tempdir().unwrap();
406        let c = Config::load(&dir.path().join("nope.toml")).unwrap();
407        assert_eq!(c.network.relay_mode, "default");
408    }
409
410    #[test]
411    fn roster_url_and_poll_interval_parse_with_defaults() {
412        // No [roster] table → url None, poll 1h default.
413        let c = Config::from_toml_str("").unwrap();
414        assert!(c.roster.url.is_none());
415        assert_eq!(c.roster.poll_interval_seconds(), 3600);
416        // A configured url + poll interval.
417        let c = Config::from_toml_str(
418            "[roster]\nurl = \"https://intranet.acme.com/roster.json\"\npoll_interval = \"30m\"\n",
419        )
420        .unwrap();
421        assert_eq!(
422            c.roster.url.as_deref(),
423            Some("https://intranet.acme.com/roster.json")
424        );
425        assert_eq!(c.roster.poll_interval_seconds(), 30 * 60);
426        // An unparseable poll_interval falls back to the hourly default (never disables the poll).
427        let c = Config::from_toml_str("[roster]\npoll_interval = \"never\"\n").unwrap();
428        assert_eq!(c.roster.poll_interval_seconds(), 3600);
429        // The url is additive: setting only grace_period keeps url None + the default poll.
430        let c = Config::from_toml_str("[roster]\ngrace_period = \"24h\"\n").unwrap();
431        assert!(c.roster.url.is_none());
432        assert_eq!(c.roster.poll_interval_seconds(), 3600);
433    }
434
435    #[test]
436    fn roster_max_staleness_defaults_to_24h_and_parses() {
437        // No [roster] table → the 24h freshness bound (the default).
438        let c = Config::from_toml_str("").unwrap();
439        assert_eq!(c.roster.max_staleness_seconds(), 24 * 3600);
440        // A configured value parses (units, like grace_period).
441        let c = Config::from_toml_str("[roster]\nmax_staleness = \"6h\"\n").unwrap();
442        assert_eq!(c.roster.max_staleness_seconds(), 6 * 3600);
443        // An unparseable value falls back to the 24h default (never disables the freshness bound).
444        let c = Config::from_toml_str("[roster]\nmax_staleness = \"forever\"\n").unwrap();
445        assert_eq!(c.roster.max_staleness_seconds(), 24 * 3600);
446        // Additive: setting only grace_period keeps the 24h max_staleness default.
447        let c = Config::from_toml_str("[roster]\ngrace_period = \"48h\"\n").unwrap();
448        assert_eq!(c.roster.max_staleness_seconds(), 24 * 3600);
449    }
450
451    #[test]
452    fn roster_grace_defaults_to_72h_and_parses_units() {
453        // Absent `[roster]` → the 72h default.
454        let c = Config::from_toml_str("").unwrap();
455        assert_eq!(c.roster.grace_seconds(), 72 * 3600);
456        // Hours / days / minutes / seconds / bare-seconds all resolve to seconds.
457        for (body, want) in [
458            ("[roster]\ngrace_period = \"24h\"\n", 24 * 3600),
459            ("[roster]\ngrace_period = \"72h\"\n", 72 * 3600),
460            ("[roster]\ngrace_period = \"1d\"\n", 24 * 3600),
461            ("[roster]\ngrace_period = \"30m\"\n", 30 * 60),
462            ("[roster]\ngrace_period = \"90s\"\n", 90),
463            ("[roster]\ngrace_period = \"3600\"\n", 3600), // bare seconds
464        ] {
465            assert_eq!(
466                Config::from_toml_str(body).unwrap().roster.grace_seconds(),
467                want,
468                "{body}"
469            );
470        }
471    }
472
473    #[test]
474    fn roster_grace_unparseable_or_negative_falls_back_to_default() {
475        // A garbage / negative / overflowing grace never disables degraded serving — it defaults.
476        for body in [
477            "[roster]\ngrace_period = \"seventy-two hours\"\n",
478            "[roster]\ngrace_period = \"-5h\"\n",
479            "[roster]\ngrace_period = \"18446744073709551615d\"\n", // overflows the checked_mul
480            "[roster]\ngrace_period = \"\"\n",
481        ] {
482            assert_eq!(
483                Config::from_toml_str(body).unwrap().roster.grace_seconds(),
484                72 * 3600,
485                "{body}"
486            );
487        }
488    }
489
490    #[test]
491    fn services_parse_run_and_socket() {
492        let c = Config::from_toml_str(concat!(
493            "[services.notes]\nrun = [\"npx\", \"server\"]\nallow = [\"bob\"]\n",
494            "[services.kb]\nsocket = \"/run/kb.sock\"\nallow = [\"team-eng\"]\n",
495        ))
496        .unwrap();
497        let notes = c.services.get("notes").unwrap();
498        assert!(
499            matches!(notes.backend_result(), Ok(Backend::Run(cmd)) if cmd == &["npx".to_string(), "server".to_string()][..])
500        );
501        assert_eq!(notes.allow, vec!["bob".to_string()]);
502        assert!(
503            matches!(c.services.get("kb").unwrap().backend_result(), Ok(Backend::Socket(p)) if p == "/run/kb.sock")
504        );
505    }
506
507    #[test]
508    fn service_with_both_run_and_socket_is_an_error() {
509        let e = Config::from_toml_str("[services.x]\nrun=[\"a\"]\nsocket=\"/s\"\nallow=[]\n");
510        // exactly one backend kind is required — validate at access time.
511        assert!(
512            e.unwrap()
513                .services
514                .get("x")
515                .unwrap()
516                .backend_result()
517                .is_err()
518        );
519    }
520
521    #[test]
522    fn identity_reads_user_id_and_user_key() {
523        let toml = "[identity]\n\
524            org_id = \"acme\"\n\
525            org_root_pk = \"b64u:AAAA\"\n\
526            user_id = \"alice\"\n\
527            user_key = \"/home/alice/.config/mcpmesh/user.key\"\n";
528        let cfg: Config = toml::from_str(toml).unwrap();
529        assert_eq!(cfg.identity.user_id.as_deref(), Some("alice"));
530        assert_eq!(
531            cfg.identity.user_key.as_deref(),
532            Some(std::path::Path::new("/home/alice/.config/mcpmesh/user.key"))
533        );
534        // Absent → None (pure-pairing / operator-only node).
535        let bare: Config = toml::from_str("[identity]\n").unwrap();
536        assert!(bare.identity.user_id.is_none() && bare.identity.user_key.is_none());
537    }
538}