1
2
3
4
5
6
7
8
9
10
11
12
13
14
15
16
17
18
19
20
21
22
23
24
25
26
27
28
29
30
31
32
33
34
35
36
37
38
39
//! Portable SIMD kernels using macerator.
//!
//! Replaces platform-specific implementations (neon.rs) with a single
//! portable implementation that auto-dispatches to NEON/AVX2/SSE/SIMD128/scalar.
use ;
use *;
/// Threshold for parallel execution (elements).
/// For memory-bound operations, parallelism helps when data exceeds L3 cache.
const PARALLEL_THRESHOLD: usize = 4 * 1024 * 1024;
const CHUNK_SIZE: usize = 4096;
pub use *;
pub use *;
pub use *;
pub use *;
pub use *;
// ============================================================================
// Tests
// ============================================================================