Skip to main content

rucc_codegen/select/
x86_64.rs

1//! The x86-64 lowering table.
2//!
3//! Everything below the module comment is generated from `rules/x86-64.rules` by `rucc-rules`
4//! when this crate is built, and none of it is in the repository. The rule file is the only
5//! place the rules are written, which is what makes the table that is matched with and the
6//! table `rucc-verify` proves things about the same table.
7//!
8//! To read the rules, read the rule file. To read the automaton they compile into, build the
9//! crate and read `x86-64.rs` under the build directory, which is a file worth looking at once
10//! for the shape of it and never again.
11
12// A guard is emitted as the comparison the rule writes, so a rule saying a shift count is at
13// least zero and less than the width comes out as two comparisons rather than as a range. That
14// is deliberate: the generated line and the rule it came from should read the same, and the
15// suggestion to write it another way is advice for somebody editing code, which nobody here is.
16#![allow(clippy::manual_range_contains)]
17
18include!(concat!(env!("OUT_DIR"), "/x86-64.rs"));
19
20#[cfg(test)]
21mod tests {
22    use rucc_target::x86_64;
23
24    use super::TABLE;
25    use crate::select::Piece;
26
27    /// The prefix a rule file puts in front of a machine term, which is how it says which target
28    /// the term belongs to. It is not part of the opcode.
29    const PREFIX: &str = "x64.";
30
31    /// The two address constructors, which are not instructions. An addressing mode is an
32    /// argument to `lea` and to every memory operand after it, so it is written as a term in the
33    /// rule file and built by the selector into the instruction that takes it.
34    const AMODES: &[&str] =
35        &["amode_base_index_scale", "amode_index_scale", "amode_base", "amode_base_offset"];
36
37    /// Every head this table can write, in and under the replacements.
38    fn heads() -> Vec<&'static str> {
39        let mut found: Vec<&'static str> = TABLE
40            .rules
41            .iter()
42            .flat_map(|rule| rule.replacement.iter())
43            .filter_map(|piece| match piece {
44                Piece::App { head, .. } => Some(*head),
45                _ => None,
46            })
47            .collect();
48        found.sort_unstable();
49        found.dedup();
50        found
51    }
52
53    #[test]
54    fn every_instruction_the_table_writes_is_described() {
55        for head in heads() {
56            if AMODES.contains(&head) {
57                continue;
58            }
59            let opcode = head.strip_prefix(PREFIX).unwrap_or_else(|| {
60                panic!("{head} is neither an x86-64 term nor an addressing mode")
61            });
62            assert!(
63                x86_64::form(opcode).is_some(),
64                "{head} is selected by a rule and `rucc_target::x86_64` does not say what it \
65                 does with its operands"
66            );
67        }
68    }
69
70    /// The order the operands of a store are written in, which is the IR's and not a choice this
71    /// file makes.
72    ///
73    /// A pattern is matched against an instruction's operand list by position, so a rule that
74    /// names the address where the IR holds the value is a rule that stores to the value and
75    /// writes the address into memory. Nothing in a proof would catch it, because a proof is
76    /// about the rule file agreeing with itself, and both halves would be wrong in the same way.
77    /// `rucc_ir::Builder::store` takes the value first and the machine instruction takes it last,
78    /// which is why the two halves of one of these rules read in opposite orders.
79    #[test]
80    fn a_store_is_written_with_the_value_first_because_that_is_where_the_ir_keeps_it() {
81        let mut seen = 0;
82        for rule in TABLE.rules {
83            let Some(rest) = rule.pattern.strip_prefix("(store.") else { continue };
84            let (width, operands) = rest.split_once(' ').expect("a store takes operands");
85            assert!(
86                operands.starts_with(&format!("(value.{width} ")),
87                "line {}: {} binds something other than the value it is storing first",
88                rule.line,
89                rule.pattern
90            );
91            assert!(
92                operands.contains("(value.i64 "),
93                "line {}: {} reaches no address",
94                rule.line,
95                rule.pattern
96            );
97            seen += 1;
98        }
99        assert_eq!(seen, 16, "the store rules moved and this test did not follow them");
100    }
101
102    /// Every comparison can be made against a constant as well as against a register.
103    ///
104    /// Four comparisons in five in the corpus are against a constant, and without a rule for one
105    /// the constant is loaded into a register first, which is an instruction and a register the
106    /// machine never needed. A missing width or a missing condition would not fail anything else:
107    /// the register rule still matches, the output is still correct, and the only sign is code
108    /// that is one instruction longer in a place nobody is looking. So the two lists are counted
109    /// against each other here.
110    ///
111    /// What this cannot check is that the condition on the immediate rule is the right one, since
112    /// both halves of a wrong pair would be a consistent pair. That is what the `spec` clause is
113    /// for, and `rucc-verify` is what reads it.
114    #[test]
115    fn a_comparison_against_a_constant_is_written_for_every_one_against_a_register() {
116        let mut against_register = Vec::new();
117        let mut against_constant = Vec::new();
118        for rule in TABLE.rules {
119            let Some(rest) = rule.pattern.strip_prefix("(icmp_") else { continue };
120            let (condition, operands) = rest.split_once(".i1 ").expect("a comparison takes two");
121            let width = operands
122                .strip_prefix("(value.")
123                .and_then(|rest| rest.split_once(' '))
124                .map(|(width, _)| width)
125                .expect("a comparison reads a value first");
126            let named = format!("{condition}.{width}");
127            if operands.contains("(iconst.") {
128                // The constant is the second operand and never the first, because a comparison is
129                // not symmetric and the same condition on the other side means the opposite.
130                assert!(
131                    !operands.starts_with("(iconst."),
132                    "line {}: {} compares a constant against a value",
133                    rule.line,
134                    rule.pattern
135                );
136                against_constant.push(named);
137            } else {
138                against_register.push(named);
139            }
140        }
141        against_register.sort_unstable();
142        against_constant.sort_unstable();
143        assert_eq!(against_register, against_constant);
144        assert_eq!(against_register.len(), 40, "ten conditions at four widths");
145    }
146
147    /// The instructions the calling convention writes rather than a rule.
148    ///
149    /// Three kinds of them. Naming the register an argument arrived in, where an argument is
150    /// depends on its position in the signature and on the classification of every argument before
151    /// it, and a rule pattern sees one term and has no way to say any of that, so `crate::abi`
152    /// builds these from the convention instead. Calling a name is the same the other way round:
153    /// what its operands are is whatever the signature made them, and a call through an address is
154    /// the same instruction with one operand more.
155    ///
156    /// The second half of a value that comes back in two registers is the third. A return of one
157    /// value is a rule, because where that value goes depends on nothing but the value, which is
158    /// exactly what a rule can say. A return of two is not, because which register the second half
159    /// is in depends on the first half: the two register files are counted separately, so a
160    /// `double` and a `long` both come back at place zero and two `long`s do not.
161    const CONVENTION: &[&str] = &[
162        "arg_val_8",
163        "arg_val_16",
164        "arg_val_32",
165        "arg_val_64",
166        "arg_val_f32",
167        "arg_val_f64",
168        "arg_val_f128",
169        "ret_val2_8",
170        "ret_val2_16",
171        "ret_val2_32",
172        "ret_val2_64",
173        "ret_val2_f32",
174        "ret_val2_f64",
175        "ret_val2_f128",
176        "call",
177        "call_reg",
178    ];
179
180    /// The instructions the block layout writes rather than a rule.
181    ///
182    /// A rule sees one branch and the layout is about the order of every block in the function, so
183    /// which arm falls through is not something any pattern could say. That answer is what decides
184    /// whether the jump goes to the arm the condition is true for or the other one, and whether
185    /// there is a second jump after it, so all of these are written where the answer is.
186    ///
187    /// The comparisons are here for a second reason on top of that one. A branch on a comparison
188    /// is a comparison and a jump on the flags it set, and the flags are not a value: no pattern
189    /// could bind one and no `spec` clause could say anything about one. So the pair is put
190    /// together by the layout, out of a comparison a rule did select and the branch behind it,
191    /// which is the same argument `rucc_target::x86_64::Form::CmpSet` is one form rather than two
192    /// under.
193    const LAYOUT: &[&str] = &[
194        "test_rr_8",
195        "cmp_rr_8",
196        "cmp_rr_16",
197        "cmp_rr_32",
198        "cmp_rr_64",
199        "cmp_ri_8",
200        "cmp_ri_16",
201        "cmp_ri_32",
202        "cmp_ri_64",
203        "jcc_e",
204        "jcc_ne",
205        "jcc_l",
206        "jcc_le",
207        "jcc_g",
208        "jcc_ge",
209        "jcc_b",
210        "jcc_be",
211        "jcc_a",
212        "jcc_ae",
213        "jmp",
214    ];
215
216    /// The instructions the compare pass writes rather than a rule.
217    ///
218    /// The other half of the argument the comparisons above are here under. A rule selects a
219    /// comparison that keeps its answer in a byte, because that is the shape a value has. What is
220    /// left of one when the machine has already made the comparison is the byte with no comparison
221    /// in front of it, and there is no pattern for that: the term it would compute is the same term
222    /// the full comparison computes, and what makes the short one right is the instruction three
223    /// places back rather than anything about the value. So `crate::compare` writes them by name,
224    /// in place of a comparison it found was already made.
225    const COMPARE: &[&str] = &[
226        "set_e", "set_ne", "set_l", "set_le", "set_g", "set_ge", "set_b", "set_be", "set_a",
227        "set_ae",
228    ];
229
230    /// The instruction a computed `goto` is written as rather than a rule.
231    ///
232    /// The one branch `crate::lower` writes by name, and the one the block layout does not write
233    /// either. What it reads is the address, which a pattern could have bound, so it is not
234    /// exempt for the reason the branches above are. What no pattern can say is the rest of it:
235    /// how many arms the block has, which is every label of the function the program took the
236    /// address of, and a rule says what an instruction reads rather than where a block goes.
237    const LABELS: &[&str] = &["jmp_reg"];
238
239    /// The instruction the memory model writes rather than a rule.
240    ///
241    /// A barrier computes nothing, so there is no equality for the solver to discharge and no
242    /// pattern for a rule to be written as. What makes it the right answer is what the machine
243    /// promises about the order two other instructions become visible in, which is a claim about
244    /// the program around it rather than about any value. `crate::lower` writes it by name, at the
245    /// strongest ordering and nowhere else, and `crate::expand` says why the strongest is the only
246    /// one that costs anything here.
247    const BARRIER: &[&str] = &["mfence"];
248
249    /// The instruction a program stops on, which `crate::lower` writes rather than a rule.
250    ///
251    /// The first half of the barrier's reason and not the second. It computes nothing, so there is
252    /// no equality for the solver and no pattern for a rule. What makes it right is not a claim
253    /// about the order anything becomes visible in either: it is what the operating system does
254    /// with the fault, which is a fact about neither the values nor the program around it.
255    const STOP: &[&str] = &["ud2"];
256
257    /// The instructions that are a hint rather than a computation.
258    ///
259    /// The same shape of exemption the barrier gets and for a reason one step further out. A
260    /// barrier computes nothing and still has to be where it is, so there is at least a claim about
261    /// the program around it. A prefetch does not even have that: a machine that drops the whole
262    /// instruction runs the program correctly, because the only thing it can change is how long the
263    /// program takes.
264    ///
265    /// So there is no equality for the solver and no pattern for a rule, and which of the four a
266    /// program gets is decided by a number in the builtin's own arguments rather than by anything
267    /// about the value being prefetched. `crate::lower` writes them by name, out of the hint the IR
268    /// carries beside the instruction.
269    const HINT: &[&str] = &["prefetch_nta", "prefetch_t0", "prefetch_t1", "prefetch_t2"];
270
271    /// The instructions nothing but an `asm` statement asks for.
272    ///
273    /// One step further out again. A prefetch is a hint and is still something the compiler decides
274    /// to write, out of a builtin the program called. This is a hint the program wrote down as an
275    /// instruction, by name, in a template, and nothing else in the language reaches it: there is no
276    /// builtin for it, no rule could match a term that produces it because it produces no value, and
277    /// `crate::lower` writes it only because [`rucc_target::x86_64::read`] found the name in a
278    /// template and said which opcode that is.
279    const TEMPLATE: &[&str] = &["pause"];
280
281    /// The instructions that produce two values, which is one more than a rule can name.
282    ///
283    /// A rule replaces a term with a term, and a term is the value one instruction computes. A
284    /// compare and exchange computes two: what it found at the address, and whether what it found
285    /// was what the program expected. There is no way to write the second one down in the rule
286    /// language, and inventing one would be inventing a language for a single instruction.
287    ///
288    /// So `crate::lower` writes it by name, the way it writes the barrier by name, and for a reason
289    /// that is about the rule language rather than about the machine. What the solver would have
290    /// been asked to prove about it is the easy half in any case: the arithmetic is a comparison
291    /// and a select, and what is hard is that the whole of it happens at once, which is the same
292    /// claim about the program around it that a barrier makes.
293    const ATOMIC: &[&str] = &["cmpxchg_8", "cmpxchg_16", "cmpxchg_32", "cmpxchg_64"];
294
295    /// The instructions whose operation is in the payload rather than in the head.
296    ///
297    /// A different exemption from the one above, on instructions that produce one value each and so
298    /// could be named by a rule if the rule had anything to match on. The head a pattern matches is
299    /// an opcode and a type, and every read modify write in the IR is the one opcode `atomic_rmw`.
300    /// Which of the thirteen operations it performs is carried beside the instruction rather than in
301    /// its name, so a pattern written for the exchange would match the add and the nand as well, and
302    /// the rule language has no way to look at what a rule matched to tell them apart.
303    ///
304    /// Giving each operation its own opcode is the other way out and is a worse trade: it is
305    /// thirteen opcodes at four widths where the IR wants one, and every pass that treats a read
306    /// modify write as one thing would then have a list of fifty two.
307    ///
308    /// So `crate::lower` writes these by name too. Three operations here, out of the thirteen: the
309    /// bitwise ones need a loop around a compare and exchange, which is control flow and so is built
310    /// before selection rather than during it, and they are the rest of `tamnd/rucc#311`.
311    const PAYLOAD: &[&str] =
312        &["xchg_8", "xchg_16", "xchg_32", "xchg_64", "xadd_8", "xadd_16", "xadd_32", "xadd_64"];
313
314    /// The instructions a frame writes rather than a rule.
315    ///
316    /// A prologue, an epilogue, a copy, a spill and a reload are not in the program. They are what
317    /// the allocator's answer costs, so they are written after it, by `crate::finish` reading
318    /// `x86_64::FRAME`. Six of the names that describes are already reachable from a rule, since a
319    /// prologue taking its frame is a subtraction and a spill is a store, and those are not here:
320    /// this is only the ones nothing else can reach.
321    const FRAME: &[&str] = &[
322        "push_64",
323        "pop_64",
324        "ret",
325        "mov_rr_64",
326        "movaps_rr",
327        // The touch a probing prologue puts on each page as it reaches it, the landing pad a
328        // prologue opens with, and the byte that does nothing which one reserves room with. All
329        // three are written by a frame and none on a command line that did not ask for it.
330        "or_mi_8",
331        "endbr64",
332        "nop",
333    ];
334
335    /// The instructions that reach the x87 stack, which are selected but not from here.
336    ///
337    /// A third kind of exemption, and the same reason all the way down the list.
338    ///
339    /// Every one of these is written by `crate::lower`, as part of a group rather than on its own.
340    /// What one of them leaves behind and the next picks up is the top of the x87 stack, which is
341    /// not a register anything allocates from and not a value a pattern could bind, so a rule
342    /// could neither match the middle of a group nor name what its replacement produced. And an
343    /// add here reads two addresses and writes a third, where one machine IR instruction carries
344    /// one addressing mode, so the group cannot be folded into a single opcode the way
345    /// `ucomisd_set_e` folds a comparison and a `setcc` either.
346    ///
347    /// So these are exempt for the reason `FRAME` is exempt rather than for the reason the list
348    /// below is, and they will stay exempt. Two of them are not reached by anything yet all the
349    /// same: `fsub_p` and `fdiv_p` are the other direction of the subtraction and the division,
350    /// which a code generator that pushed its operands the other way round would need and this one
351    /// does not. `fabs` is a third, since C spells that as a call to a library function.
352    const X87: &[&str] = &[
353        "fld_t",
354        "fstp_t",
355        "fld_s",
356        "fld_l",
357        "fild_l",
358        "fild_ll",
359        "fstp_s",
360        "fstp_l",
361        "fistp_l",
362        "fistp_ll",
363        "fnstcw",
364        "fldcw",
365        "fadd_p",
366        "fsub_p",
367        "fsubr_p",
368        "fmul_p",
369        "fdiv_p",
370        "fdivr_p",
371        "fchs",
372        "fabs",
373        "fucomip_set_a",
374        "fucomip_set_ae",
375        "fucomip_set_b",
376        "fucomip_set_be",
377        "fucomip_set_e",
378        "fucomip_set_ne",
379        "fucomip_set_p",
380        "fucomip_set_np",
381        "fucomip_set_e_and_np",
382        "fucomip_set_ne_or_p",
383    ];
384
385    /// The instructions no rule selects yet, because the rules that selected them were taken out.
386    ///
387    /// A different kind of exemption from the three above. Those say an instruction is written
388    /// somewhere a rule cannot reach and always will be. These say nobody reaches one at all right
389    /// now, and name the work that puts the rules back.
390    ///
391    /// The rules went out under `tamnd/rucc#368`. C promotes the operands of an arithmetic
392    /// operator to `int`, so a byte add and a two byte compare are things no C program asks the
393    /// back end for, and the rules at those widths sat proved and never selected over the whole
394    /// torture corpus at every optimization level. The width narrowing pass in `tamnd/rucc#375` is
395    /// what asks for them, and the rules come back with it.
396    ///
397    /// The descriptions stayed. A description says what an x86-64 instruction is, how long it is
398    /// and how it encodes, and that is true whether or not anything selects it. Taking them out
399    /// would be deleting a correct account of the machine to make a list shorter, and putting them
400    /// back is then a second thing to get right rather than a line of a rule file.
401    const NARROW: &[&str] = &[
402        // Three of the two address forms against an immediate. The `narrow` pass does write the
403        // shape, since `char c = a | 1;` narrows to a byte `or` against a byte constant, and no
404        // rule selects these yet: the constant goes into a register and the register with
405        // register rule takes it. Their `add`, `sub` and `and` siblings do have rules and are
406        // reached by the bitfield lowering, so this is six rules missing rather than a shape
407        // nothing writes.
408        "or_ri_8",
409        "or_ri_16",
410        "xor_ri_8",
411        "xor_ri_16",
412        "imul_ri_8",
413        "imul_ri_16",
414        // The divides, which are four instructions per width because the quotient and the
415        // remainder come out of one division in two different registers. `narrow` refuses these
416        // on purpose: the most negative byte over minus one is a defined hundred and twenty eight
417        // at four bytes and is the overflow that raises at one, so narrowing a division wants a
418        // range that rules the pair out and there is no range analysis yet.
419        "idiv_quo_8",
420        "idiv_quo_16",
421        "idiv_rem_8",
422        "idiv_rem_16",
423        "div_quo_8",
424        "div_quo_16",
425        "div_rem_8",
426        "div_rem_16",
427        // The shifts by a value, whose count is in `cl` whatever the width being shifted is. The
428        // same refusal for the same kind of reason: a count of twenty is a defined shift to zero
429        // at four bytes and is poison at one, so only a count that is a constant below the narrow
430        // width narrows, and that one selects the immediate forms which do have rules.
431        "shl_rcl_8",
432        "shl_rcl_16",
433        "shr_rcl_8",
434        "shr_rcl_16",
435        "sar_rcl_8",
436        "sar_rcl_16",
437    ];
438
439    #[test]
440    fn every_instruction_exempt_from_a_rule_is_one_a_frame_really_writes() {
441        // The same claim as the one about the convention, so that this list cannot grow an opcode
442        // that no frame asks for. In the order `x86_64::FRAME` names them, the copies after the
443        // return because there is one set of them per class the allocator may spill.
444        let frame = &x86_64::FRAME;
445        let mut written = vec![frame.push, frame.pop, frame.ret];
446        for class in frame.classes {
447            written.extend([class.mov, class.load, class.store]);
448        }
449        // And the touch a probing prologue puts on a page, which the target names as an option
450        // because a target with no instruction that writes an address without changing it takes
451        // every frame in one subtraction and has nothing to exempt.
452        written.extend(frame.probe.map(|probe| probe.inst));
453        // And the landing pad and the byte that does nothing, which are options for the same
454        // reason.
455        written.extend(frame.landing);
456        written.extend(frame.pad);
457        // What is left after the ones a rule already reaches, which are the loads and the stores of
458        // both register files, since those are the same instructions a program's own reads and
459        // writes of memory are. The vector pair joined them with the rules for a quad float, and a
460        // spill of one is now the same instruction as a program reading a `_Float128` variable.
461        written.retain(|opcode| !heads().contains(&format!("{PREFIX}{opcode}").as_str()));
462        assert_eq!(written, FRAME);
463    }
464
465    #[test]
466    fn every_instruction_exempt_from_a_rule_is_one_the_convention_really_writes() {
467        // An exemption list that nothing checks is a hole, since an opcode dropped into it stops
468        // being covered by either direction of the pinning. These are the ones `crate::abi` can
469        // name, at the four integer widths and the two float formats it has names for an
470        // argument in, and no others.
471        let strip = |head: &'static str| head.strip_prefix(PREFIX).expect("an x86-64 term");
472        let named = |ty| strip(crate::abi::head_of(ty).expect("every width the pseudos cover"));
473        // The second half of a pair at place one, which is the place a rule cannot name. The first
474        // half at place zero is `ret_val_*` and is reached by a rule, so it is not on this list.
475        let second = |ty| strip(crate::abi::ret_of(ty, 1).expect("every width the pseudos cover"));
476        let widths = || {
477            [8, 16, 32, 64].into_iter().map(rucc_ir::Type::int).chain(
478                [rucc_ir::Float::F32, rucc_ir::Float::F64, rucc_ir::Float::F128]
479                    .map(rucc_ir::Type::float),
480            )
481        };
482        let written: Vec<&str> = widths()
483            .map(named)
484            .chain(widths().map(second))
485            .chain([strip(crate::abi::CALL), strip(crate::abi::CALL_REG)])
486            .collect();
487        assert_eq!(written, CONVENTION);
488    }
489
490    /// The same claim about the block layout's list, which is longer than it looks.
491    ///
492    /// A name here that the layout does not write is an opcode exempted from needing a rule and
493    /// reached by nothing, and a name the layout writes that is not here is a failing test in
494    /// `every_described_instruction_is_reachable_from_a_rule` with a misleading message. Both are
495    /// avoided by taking the list from `rucc_target::x86_64::BRANCH` rather than believing it.
496    #[test]
497    fn every_instruction_exempt_from_a_rule_is_one_the_block_layout_really_writes() {
498        let branch = &x86_64::BRANCH;
499        // Eighty entries name sixteen instructions between them, so this is a set rather than a
500        // list and both sides are sorted before they are held against each other. What the order
501        // of the list itself is for is reading it.
502        let mut written: Vec<&str> = vec![branch.test, branch.jump];
503        written.extend(branch.fused.iter().map(|fusion| fusion.cmp));
504        written.extend(branch.fused.iter().flat_map(|fusion| [fusion.if_true, fusion.if_false]));
505        written.sort_unstable();
506        written.dedup();
507        let mut exempt = LAYOUT.to_vec();
508        exempt.sort_unstable();
509        assert_eq!(written, exempt);
510    }
511
512    /// The same claim about the compare pass. What it writes is what the flag description says is
513    /// left of a comparison, so the exemption is taken from that rather than typed out twice, and
514    /// an entry added there without a rule to go with it shows up here rather than in a build that
515    /// fails somewhere else.
516    #[test]
517    fn every_instruction_exempt_from_a_rule_is_one_the_compare_pass_really_writes() {
518        let mut written: Vec<&str> =
519            x86_64::FLAGS.compares.iter().filter_map(|entry| entry.kept).collect();
520        written.sort_unstable();
521        written.dedup();
522        let mut exempt = COMPARE.to_vec();
523        exempt.sort_unstable();
524        assert_eq!(written, exempt);
525    }
526
527    /// And the same claim about the one the lowering writes, held against the name the target gave
528    /// it rather than against the spelling written above.
529    #[test]
530    fn the_instruction_a_computed_goto_is_exempt_for_is_the_one_the_target_names() {
531        assert_eq!(LABELS, [x86_64::BRANCH.indirect]);
532    }
533
534    #[test]
535    fn every_described_instruction_is_reachable_from_a_rule() {
536        let written = heads();
537        for &(opcode, _) in x86_64::INSTS {
538            if CONVENTION.contains(&opcode) || LAYOUT.contains(&opcode) || FRAME.contains(&opcode) {
539                continue;
540            }
541            if NARROW.contains(&opcode) || BARRIER.contains(&opcode) || X87.contains(&opcode) {
542                continue;
543            }
544            if ATOMIC.contains(&opcode) || PAYLOAD.contains(&opcode) || HINT.contains(&opcode) {
545                continue;
546            }
547            if COMPARE.contains(&opcode) || TEMPLATE.contains(&opcode) {
548                continue;
549            }
550            if LABELS.contains(&opcode) || STOP.contains(&opcode) {
551                continue;
552            }
553            let head = format!("{PREFIX}{opcode}");
554            assert!(
555                written.contains(&head.as_str()),
556                "{opcode} is described and no rule in {} selects it",
557                TABLE.source
558            );
559        }
560    }
561
562    /// The same claim about the barrier as the ones above make about the convention and the frame:
563    /// the list holds instructions this target really describes, and holds only the ones that have
564    /// no operands, since an instruction with an operand is one a rule could have been written for.
565    #[test]
566    fn every_instruction_exempt_from_a_rule_is_one_the_memory_model_really_writes() {
567        for &opcode in BARRIER {
568            let form = x86_64::form(opcode).expect("an instruction this target describes");
569            assert!(form.operands().is_empty(), "{opcode} has operands, so a rule could name it");
570        }
571    }
572
573    /// The same claim about the instruction a program stops on, which is the barrier's shape
574    /// exactly: no operands, because an instruction with one is an instruction a rule could have
575    /// been written for, and no addressing mode either, because it is given nothing at all.
576    #[test]
577    fn the_instruction_exempt_from_a_rule_because_it_stops_the_program_is_bare() {
578        for &opcode in STOP {
579            let form = x86_64::form(opcode).expect("an instruction this target describes");
580            assert!(form.operands().is_empty(), "{opcode} has operands, so a rule could name it");
581            assert!(!form.takes_mem(), "{opcode} is given an address and stopping needs none");
582        }
583    }
584
585    /// The same claim about the hints, with the one difference between them written down. A hint is
586    /// given an address and nothing else, so it has no operands for the reason a barrier has none
587    /// and it does carry an addressing mode, which is what a rule would have had to match on.
588    #[test]
589    fn every_instruction_exempt_from_a_rule_because_it_is_a_hint_is_given_only_an_address() {
590        for &opcode in HINT {
591            let form = x86_64::form(opcode).expect("an instruction this target describes");
592            assert!(form.operands().is_empty(), "{opcode} has operands, so a rule could name it");
593            assert!(form.takes_mem(), "{opcode} is a hint about an address and is given none");
594        }
595    }
596
597    /// The same claim about the template list. An instruction is exempt for this reason exactly
598    /// when it computes nothing and is given nothing, since anything with an operand or an address
599    /// is something a rule or a builtin could have been written for, and the whole of the claim is
600    /// that there was nowhere else for it to come from.
601    #[test]
602    fn every_instruction_exempt_from_a_rule_because_only_a_template_asks_for_it_is_bare() {
603        for &opcode in TEMPLATE {
604            let form = x86_64::form(opcode).expect("an instruction this target describes");
605            assert!(form.operands().is_empty(), "{opcode} has operands, so a rule could name it");
606            assert!(!form.takes_mem(), "{opcode} is given an address, so a rule could name it");
607        }
608    }
609
610    /// The same claim about the atomic list, read off the thing that put the entry there: an
611    /// instruction is exempt for this reason exactly when it writes more than one value, and an
612    /// instruction that writes one is one a rule could have been written for.
613    #[test]
614    fn every_instruction_exempt_from_a_rule_is_one_that_writes_more_than_one_value() {
615        let written = heads();
616        for &opcode in ATOMIC {
617            let form = x86_64::form(opcode).expect("an instruction this target describes");
618            let writes = form.operands().iter().filter(|desc| desc.role.is_def()).count();
619            assert!(writes > 1, "{opcode} writes one value, so a rule could name it");
620            assert!(
621                !written.contains(&format!("{PREFIX}{opcode}").as_str()),
622                "a rule in {} selects {opcode}, which `crate::lower` also writes by hand",
623                TABLE.source
624            );
625        }
626    }
627
628    /// The same claim about the payload list, read off the thing that puts an entry there.
629    ///
630    /// Two halves. Each of these writes one value, which is what says the reason above is not the
631    /// reason here, so a list that grew to cover an instruction the atomic list should have had
632    /// fails. And there really is more than one operation behind the one IR opcode, which is the
633    /// whole of why a pattern cannot name any of them, and is a fact about the IR that would stop
634    /// being true if the operations were ever given opcodes of their own.
635    #[test]
636    fn every_instruction_exempt_because_its_operation_is_beside_it_writes_one_value() {
637        assert!(
638            rucc_ir::RmwOp::all().count() > 1,
639            "one operation per opcode would be a head a rule could match"
640        );
641        let written = heads();
642        for &opcode in PAYLOAD {
643            let form = x86_64::form(opcode).expect("an instruction this target describes");
644            let writes = form.operands().iter().filter(|desc| desc.role.is_def()).count();
645            assert_eq!(writes, 1, "{opcode} writes more than one value, so it is the other list's");
646            assert!(
647                !written.contains(&format!("{PREFIX}{opcode}").as_str()),
648                "a rule in {} selects {opcode}, which `crate::lower` also writes by hand",
649                TABLE.source
650            );
651        }
652    }
653
654    /// The staleness rule every list in this project is kept under, on the one list here whose
655    /// entries are meant to leave. A rule that starts selecting one of these is `tamnd/rucc#375`
656    /// arriving, and the entry goes with it. An entry naming an instruction nothing describes is a
657    /// misspelling, and it would sit here exempting nothing.
658    #[test]
659    fn an_instruction_a_rule_now_selects_is_off_the_list_of_the_ones_left_for_later() {
660        let written = heads();
661        for &opcode in NARROW {
662            let head = format!("{PREFIX}{opcode}");
663            assert!(
664                !written.contains(&head.as_str()),
665                "a rule in {} selects {opcode} now, so it is not waiting on tamnd/rucc#375",
666                TABLE.source
667            );
668            assert!(
669                x86_64::INSTS.iter().any(|&(described, _)| described == opcode),
670                "{opcode} is not an instruction anything describes"
671            );
672        }
673    }
674
675    /// The same staleness rule on the x87 pair, and one thing more that is particular to them.
676    ///
677    /// They are a pair. An instruction that pushes onto the x87 stack and nothing that pops off it
678    /// again would leave the stack one deeper than the function found it, which is not a mistake
679    /// the allocator or the block layout could catch, since neither of them knows the stack is
680    /// there. So the two arrive together and leave together, and that is what this says.
681    #[test]
682    fn the_x87_stack_is_reached_by_a_pair_and_by_nothing_else() {
683        let written = heads();
684        for &opcode in X87 {
685            assert!(
686                x86_64::INSTS.iter().any(|&(described, _)| described == opcode),
687                "{opcode} is not an instruction anything describes"
688            );
689            assert!(
690                !written.contains(&format!("{PREFIX}{opcode}").as_str()),
691                "a rule in {} selects {opcode}, which `crate::lower` also writes by hand",
692                TABLE.source
693            );
694        }
695        // One way onto the stack per format a value can be read from, one way off it per format a
696        // value can be written to, the control word pair that is neither, and the arithmetic. The
697        // count is here as well as in the target description because this list is what says none
698        // of them is reachable, and a name that arrived here without its partner would be a format
699        // this target can convert in one direction and not the other.
700        assert_eq!(X87.len(), 30, "twelve that move a value and eighteen that work on one");
701    }
702}