Add caching to inversion list searches

[perl5.git] / regcomp.c
diff --git a/regcomp.c b/regcomp.c

index 11f7f1d..28f2ecb 100644 (file)
--- a/regcomp.c
+++ b/regcomp.c
@@ -88,6 +88,7 @@ extern const struct regexp_engine my_reg_engine;
  
  #include "dquote_static.c"
  #include "charclass_invlists.h"
+#include "inline_invlist.c"
  
  #define HAS_NONLATIN1_FOLD_CLOSURE(i) _HAS_NONLATIN1_FOLD_CLOSURE_ONLY_FOR_USE_BY_REGCOMP_DOT_C_AND_REGEXEC_DOT_C(i)
  #define IS_NON_FINAL_FOLD(c) _IS_NON_FINAL_FOLD_ONLY_FOR_USE_BY_REGCOMP_DOT_C(c)
@@ -2661,7 +2662,7 @@ S_make_trie_failtable(pTHX_ RExC_state_t *pRExC_state, regnode *source,  regnode
   *      'ss' or not is not knowable at compile time.  It will match iff the
   *      target string is in UTF-8, unlike the EXACTFU nodes, where it always
   *      matches; and the EXACTFL and EXACTFA nodes where it never does.  Thus
- *      it can't be folded to "ss" at compile time, unlike EXACTFU does as
+ *      it can't be folded to "ss" at compile time, unlike EXACTFU does (as
   *      described in item 3).  An assumption that the optimizer part of
   *      regexec.c (probably unwittingly) makes is that a character in the
   *      pattern corresponds to at most a single character in the target string.
@@ -3663,7 +3664,7 @@ S_study_chunk(pTHX_ RExC_state_t *pRExC_state, regnode **scanp,
                 uc = utf8_to_uvchr_buf(s, s + l, NULL);
                 l = utf8_length(s, s + l);
             }
-           else if (has_exactf_sharp_s) {
+           if (has_exactf_sharp_s) {
                 RExC_seen |= REG_SEEN_EXACTF_SHARP_S;
             }
             min += l - min_subtract;
@@ -3960,6 +3961,7 @@ S_study_chunk(pTHX_ RExC_state_t *pRExC_state, regnode **scanp,
                       && !(data->flags & SF_HAS_EVAL)
                       && !deltanext     /* atom is fixed width */
                       && minnext != 0   /* CURLYM can't handle zero width */
+                      && ! (RExC_seen & REG_SEEN_EXACTF_SHARP_S) /* Nor \xDF */
                 ) {
                     /* XXXX How to optimize if data == 0? */
                     /* Optimize to a simpler form.  */
@@ -5206,6 +5208,50 @@ S_compile_runtime_code(pTHX_ RExC_state_t * const pRExC_state,
  }
  
  
+STATIC bool
+S_setup_longest(pTHX_ RExC_state_t *pRExC_state, SV* sv_longest, SV** rx_utf8, SV** rx_substr, I32* rx_end_shift, I32 lookbehind, I32 offset, I32 *minlen, STRLEN longest_length, bool eol, bool meol)
+{
+    /* This is the common code for setting up the floating and fixed length
+     * string data extracted from Perlre_op_compile() below.  Returns a boolean
+     * as to whether succeeded or not */
+
+    I32 t,ml;
+
+    if (! (longest_length
+           || (eol /* Can't have SEOL and MULTI */
+               && (! meol || (RExC_flags & RXf_PMf_MULTILINE)))
+          )
+            /* See comments for join_exact for why REG_SEEN_EXACTF_SHARP_S */
+        || (RExC_seen & REG_SEEN_EXACTF_SHARP_S))
+    {
+        return FALSE;
+    }
+
+    /* copy the information about the longest from the reg_scan_data
+        over to the program. */
+    if (SvUTF8(sv_longest)) {
+        *rx_utf8 = sv_longest;
+        *rx_substr = NULL;
+    } else {
+        *rx_substr = sv_longest;
+        *rx_utf8 = NULL;
+    }
+    /* end_shift is how many chars that must be matched that
+        follow this item. We calculate it ahead of time as once the
+        lookbehind offset is added in we lose the ability to correctly
+        calculate it.*/
+    ml = minlen ? *(minlen) : (I32)longest_length;
+    *rx_end_shift = ml - offset
+        - longest_length + (SvTAIL(sv_longest) != 0)
+        + lookbehind;
+
+    t = (eol/* Can't have SEOL and MULTI */
+         && (! meol || (RExC_flags & RXf_PMf_MULTILINE)));
+    fbm_compile(sv_longest, t ? FBMcf_TAIL : 0);
+
+    return TRUE;
+}
+
  /*
   * Perl_re_op_compile - the perl internal RE engine's function to compile a
   * regular expression into internal code.
@@ -5258,7 +5304,7 @@ Perl_re_op_compile(pTHX_ SV ** const patternp, int pat_count,
      dVAR;
      REGEXP *rx;
      struct regexp *r;
-    register regexp_internal *ri;
+    regexp_internal *ri;
      STRLEN plen;
      char  * VOL exp;
      char* xend;
@@ -6173,105 +6219,56 @@ reStudy:
         scan_commit(pRExC_state, &data,&minlen,0);
         SvREFCNT_dec(data.last_found);
  
-        /* Note that code very similar to this but for anchored string 
-           follows immediately below, changes may need to be made to both. 
-           Be careful. 
-         */
         longest_float_length = CHR_SVLEN(data.longest_float);
-       if (longest_float_length
-           || (data.flags & SF_FL_BEFORE_EOL
-               && (!(data.flags & SF_FL_BEFORE_MEOL)
-                   || (RExC_flags & RXf_PMf_MULTILINE)))) 
-        {
-            I32 t,ml;
  
-            /* See comments for join_exact for why REG_SEEN_EXACTF_SHARP_S */
-           if ((RExC_seen & REG_SEEN_EXACTF_SHARP_S)
-               || (SvCUR(data.longest_fixed)  /* ok to leave SvCUR */
-                   && data.offset_fixed == data.offset_float_min
-                   && SvCUR(data.longest_fixed) == SvCUR(data.longest_float)))
-                   goto remove_float;          /* As in (a)+. */
-
-            /* copy the information about the longest float from the reg_scan_data
-               over to the program. */
-           if (SvUTF8(data.longest_float)) {
-               r->float_utf8 = data.longest_float;
-               r->float_substr = NULL;
-           } else {
-               r->float_substr = data.longest_float;
-               r->float_utf8 = NULL;
-           }
-           /* float_end_shift is how many chars that must be matched that 
-              follow this item. We calculate it ahead of time as once the
-              lookbehind offset is added in we lose the ability to correctly
-              calculate it.*/
-           ml = data.minlen_float ? *(data.minlen_float) 
-                                  : (I32)longest_float_length;
-           r->float_end_shift = ml - data.offset_float_min
-               - longest_float_length + (SvTAIL(data.longest_float) != 0)
-               + data.lookbehind_float;
+        if (! ((SvCUR(data.longest_fixed)  /* ok to leave SvCUR */
+                   && data.offset_fixed == data.offset_float_min
+                   && SvCUR(data.longest_fixed) == SvCUR(data.longest_float)))
+            && S_setup_longest (aTHX_ pRExC_state,
+                                    data.longest_float,
+                                    &(r->float_utf8),
+                                    &(r->float_substr),
+                                    &(r->float_end_shift),
+                                    data.lookbehind_float,
+                                    data.offset_float_min,
+                                    data.minlen_float,
+                                    longest_float_length,
+                                    data.flags & SF_FL_BEFORE_EOL,
+                                    data.flags & SF_FL_BEFORE_MEOL))
+        {
             r->float_min_offset = data.offset_float_min - data.lookbehind_float;
             r->float_max_offset = data.offset_float_max;
             if (data.offset_float_max < I32_MAX) /* Don't offset infinity */
                 r->float_max_offset -= data.lookbehind_float;
-           
-           t = (data.flags & SF_FL_BEFORE_EOL /* Can't have SEOL and MULTI */
-                      && (!(data.flags & SF_FL_BEFORE_MEOL)
-                          || (RExC_flags & RXf_PMf_MULTILINE)));
-           fbm_compile(data.longest_float, t ? FBMcf_TAIL : 0);
         }
         else {
-         remove_float:
             r->float_substr = r->float_utf8 = NULL;
             SvREFCNT_dec(data.longest_float);
             longest_float_length = 0;
         }
  
-        /* Note that code very similar to this but for floating string 
-           is immediately above, changes may need to be made to both. 
-           Be careful. 
-         */
         longest_fixed_length = CHR_SVLEN(data.longest_fixed);
  
-        /* See comments for join_exact for why REG_SEEN_EXACTF_SHARP_S */
-       if (! (RExC_seen & REG_SEEN_EXACTF_SHARP_S)
-           && (longest_fixed_length
-               || (data.flags & SF_FIX_BEFORE_EOL /* Cannot have SEOL and MULTI */
-                   && (!(data.flags & SF_FIX_BEFORE_MEOL)
-                       || (RExC_flags & RXf_PMf_MULTILINE)))) )
+        if (S_setup_longest (aTHX_ pRExC_state,
+                                data.longest_fixed,
+                                &(r->anchored_utf8),
+                                &(r->anchored_substr),
+                                &(r->anchored_end_shift),
+                                data.lookbehind_fixed,
+                                data.offset_fixed,
+                                data.minlen_fixed,
+                                longest_fixed_length,
+                                data.flags & SF_FIX_BEFORE_EOL,
+                                data.flags & SF_FIX_BEFORE_MEOL))
          {
-            I32 t,ml;
-
-            /* copy the information about the longest fixed 
-               from the reg_scan_data over to the program. */
-           if (SvUTF8(data.longest_fixed)) {
-               r->anchored_utf8 = data.longest_fixed;
-               r->anchored_substr = NULL;
-           } else {
-               r->anchored_substr = data.longest_fixed;
-               r->anchored_utf8 = NULL;
-           }
-           /* fixed_end_shift is how many chars that must be matched that 
-              follow this item. We calculate it ahead of time as once the
-              lookbehind offset is added in we lose the ability to correctly
-              calculate it.*/
-            ml = data.minlen_fixed ? *(data.minlen_fixed) 
-                                   : (I32)longest_fixed_length;
-            r->anchored_end_shift = ml - data.offset_fixed
-               - longest_fixed_length + (SvTAIL(data.longest_fixed) != 0)
-               + data.lookbehind_fixed;
             r->anchored_offset = data.offset_fixed - data.lookbehind_fixed;
-
-           t = (data.flags & SF_FIX_BEFORE_EOL /* Can't have SEOL and MULTI */
-                && (!(data.flags & SF_FIX_BEFORE_MEOL)
-                    || (RExC_flags & RXf_PMf_MULTILINE)));
-           fbm_compile(data.longest_fixed, t ? FBMcf_TAIL : 0);
         }
         else {
             r->anchored_substr = r->anchored_utf8 = NULL;
             SvREFCNT_dec(data.longest_fixed);
             longest_fixed_length = 0;
         }
+
         if (ri->regstclass
             && (OP(ri->regstclass) == REG_ANY || OP(ri->regstclass) == SANY))
             ri->regstclass = NULL;
@@ -7000,32 +6997,8 @@ S_reg_scan_name(pTHX_ RExC_state_t *pRExC_state, U32 flags)
   * Some of the methods should always be private to the implementation, and some
   * should eventually be made public */
  
-#define INVLIST_LEN_OFFSET 0   /* Number of elements in the inversion list */
-#define INVLIST_ITER_OFFSET 1  /* Current iteration position */
-
-/* This is a combination of a version and data structure type, so that one
- * being passed in can be validated to be an inversion list of the correct
- * vintage.  When the structure of the header is changed, a new random number
- * in the range 2**31-1 should be generated and the new() method changed to
- * insert that at this location.  Then, if an auxiliary program doesn't change
- * correspondingly, it will be discovered immediately */
-#define INVLIST_VERSION_ID_OFFSET 2
-#define INVLIST_VERSION_ID 1064334010
-
-/* For safety, when adding new elements, remember to #undef them at the end of
- * the inversion list code section */
+/* The header definitions are in F<inline_invlist.c> */
  
-#define INVLIST_ZERO_OFFSET 3  /* 0 or 1; must be last element in header */
-/* The UV at position ZERO contains either 0 or 1.  If 0, the inversion list
- * contains the code point U+00000, and begins here.  If 1, the inversion list
- * doesn't contain U+0000, and it begins at the next UV in the array.
- * Inverting an inversion list consists of adding or removing the 0 at the
- * beginning of it.  By reserving a space for that 0, inversion can be made
- * very fast */
-
-#define HEADER_LENGTH (INVLIST_ZERO_OFFSET + 1)
-
-/* Internally things are UVs */
  #define TO_INTERNAL_SIZE(x) ((x + HEADER_LENGTH) * sizeof(UV))
  #define FROM_INTERNAL_SIZE(x) ((x / sizeof(UV)) - HEADER_LENGTH)
  
@@ -7047,7 +7020,7 @@ S__invlist_array_init(pTHX_ SV* const invlist, const bool will_have_0)
      PERL_ARGS_ASSERT__INVLIST_ARRAY_INIT;
  
      /* Must be empty */
-    assert(! *get_invlist_len_addr(invlist));
+    assert(! *_get_invlist_len_addr(invlist));
  
      /* 1^1 = 0; 1^0 = 1 */
      *zero = 1 ^ will_have_0;
@@ -7065,7 +7038,7 @@ S_invlist_array(pTHX_ SV* const invlist)
  
      /* Must not be empty.  If these fail, you probably didn't check for <len>
       * being non-zero before trying to get the array */
-    assert(*get_invlist_len_addr(invlist));
+    assert(*_get_invlist_len_addr(invlist));
      assert(*get_invlist_zero_addr(invlist) == 0
            || *get_invlist_zero_addr(invlist) == 1);
  
@@ -7076,28 +7049,6 @@ S_invlist_array(pTHX_ SV* const invlist)
                    + *get_invlist_zero_addr(invlist));
  }
  
-PERL_STATIC_INLINE UV*
-S_get_invlist_len_addr(pTHX_ SV* invlist)
-{
-    /* Return the address of the UV that contains the current number
-     * of used elements in the inversion list */
-
-    PERL_ARGS_ASSERT_GET_INVLIST_LEN_ADDR;
-
-    return (UV *) (SvPVX(invlist) + (INVLIST_LEN_OFFSET * sizeof (UV)));
-}
-
-PERL_STATIC_INLINE UV
-S_invlist_len(pTHX_ SV* const invlist)
-{
-    /* Returns the current number of elements stored in the inversion list's
-     * array */
-
-    PERL_ARGS_ASSERT_INVLIST_LEN;
-
-    return *get_invlist_len_addr(invlist);
-}
-
  PERL_STATIC_INLINE void
  S_invlist_set_len(pTHX_ SV* const invlist, const UV len)
  {
@@ -7105,7 +7056,7 @@ S_invlist_set_len(pTHX_ SV* const invlist, const UV len)
  
      PERL_ARGS_ASSERT_INVLIST_SET_LEN;
  
-    *get_invlist_len_addr(invlist) = len;
+    *_get_invlist_len_addr(invlist) = len;
  
      assert(len <= SvLEN(invlist));
  
@@ -7125,6 +7076,39 @@ S_invlist_set_len(pTHX_ SV* const invlist, const UV len)
       * Note that when inverting, SvCUR shouldn't change */
  }
  
+PERL_STATIC_INLINE IV*
+S_get_invlist_previous_index_addr(pTHX_ SV* invlist)
+{
+    /* Return the address of the UV that is reserved to hold the cached index
+     * */
+
+    PERL_ARGS_ASSERT_GET_INVLIST_PREVIOUS_INDEX_ADDR;
+
+    return (IV *) (SvPVX(invlist) + (INVLIST_PREVIOUS_INDEX_OFFSET * sizeof (UV)));
+}
+
+PERL_STATIC_INLINE IV
+S_invlist_previous_index(pTHX_ SV* const invlist)
+{
+    /* Returns cached index of previous search */
+
+    PERL_ARGS_ASSERT_INVLIST_PREVIOUS_INDEX;
+
+    return *get_invlist_previous_index_addr(invlist);
+}
+
+PERL_STATIC_INLINE void
+S_invlist_set_previous_index(pTHX_ SV* const invlist, const IV index)
+{
+    /* Caches <index> for later retrieval */
+
+    PERL_ARGS_ASSERT_INVLIST_SET_PREVIOUS_INDEX;
+
+    assert(index == 0 || index < (int) _invlist_len(invlist));
+
+    *get_invlist_previous_index_addr(invlist) = index;
+}
+
  PERL_STATIC_INLINE UV
  S_invlist_max(pTHX_ SV* const invlist)
  {
@@ -7175,8 +7159,9 @@ Perl__new_invlist(pTHX_ IV initial_size)
       * properly */
      *get_invlist_zero_addr(new_list) = UV_MAX;
  
+    *get_invlist_previous_index_addr(new_list) = 0;
      *get_invlist_version_id_addr(new_list) = INVLIST_VERSION_ID;
-#if HEADER_LENGTH != 4
+#if HEADER_LENGTH != 5
  #   error Need to regenerate VERSION_ID by running perl -E 'say int(rand 2**31-1)', and then changing the #if to the new length
  #endif
  
@@ -7199,7 +7184,7 @@ S__new_invlist_C_array(pTHX_ UV* list)
      SvPV_set(invlist, (char *) list);
      SvLEN_set(invlist, 0);  /* Means we own the contents, and the system
                                shouldn't touch it */
-    SvCUR_set(invlist, TO_INTERNAL_SIZE(invlist_len(invlist)));
+    SvCUR_set(invlist, TO_INTERNAL_SIZE(_invlist_len(invlist)));
  
      if (*get_invlist_version_id_addr(invlist) != INVLIST_VERSION_ID) {
          Perl_croak(aTHX_ "panic: Incorrect version for previously generated inversion list");
@@ -7229,11 +7214,6 @@ S_invlist_trim(pTHX_ SV* const invlist)
      SvPV_shrink_to_cur((SV *) invlist);
  }
  
-/* An element is in an inversion list iff its index is even numbered: 0, 2, 4,
- * etc */
-#define ELEMENT_RANGE_MATCHES_INVLIST(i) (! ((i) & 1))
-#define PREV_RANGE_MATCHES_INVLIST(i) (! ELEMENT_RANGE_MATCHES_INVLIST(i))
-
  #define _invlist_union_complement_2nd(a, b, output) _invlist_union_maybe_complement_2nd(a, b, TRUE, output)
  
  STATIC void
@@ -7245,7 +7225,7 @@ S__append_range_to_invlist(pTHX_ SV* const invlist, const UV start, const UV end
  
      UV* array;
      UV max = invlist_max(invlist);
-    UV len = invlist_len(invlist);
+    UV len = _invlist_len(invlist);
  
      PERL_ARGS_ASSERT__APPEND_RANGE_TO_INVLIST;
  
@@ -7326,23 +7306,68 @@ Perl__invlist_search(pTHX_ SV* const invlist, const UV cp)
       * contains <cp> */
  
      IV low = 0;
-    IV high = invlist_len(invlist);
-    const UV * const array = invlist_array(invlist);
+    IV mid;
+    IV high = _invlist_len(invlist);
+    const IV highest_element = high - 1;
+    const UV* array;
  
      PERL_ARGS_ASSERT__INVLIST_SEARCH;
  
-    /* If list is empty or the code point is before the first element, return
-     * failure. */
-    if (high == 0 || cp < array[0]) {
+    /* If list is empty, return failure. */
+    if (high == 0) {
         return -1;
      }
  
+    /* If the code point is before the first element, return failure.  (We
+     * can't combine this with the test above, because we can't get the array
+     * unless we know the list is non-empty) */
+    array = invlist_array(invlist);
+
+    mid = invlist_previous_index(invlist);
+    assert(mid >=0 && mid <= highest_element);
+
+    /* <mid> contains the cache of the result of the previous call to this
+     * function (0 the first time).  See if this call is for the same result,
+     * or if it is for mid-1.  This is under the theory that calls to this
+     * function will often be for related code points that are near each other.
+     * And benchmarks show that caching gives better results.  We also test
+     * here if the code point is within the bounds of the list.  These tests
+     * replace others that would have had to be made anyway to make sure that
+     * the array bounds were not exceeded, and give us extra information at the
+     * same time */
+    if (cp >= array[mid]) {
+        if (cp >= array[highest_element]) {
+            return highest_element;
+        }
+
+        /* Here, array[mid] <= cp < array[highest_element].  This means that
+         * the final element is not the answer, so can exclude it; it also
+         * means that <mid> is not the final element, so can refer to 'mid + 1'
+         * safely */
+        if (cp < array[mid + 1]) {
+            return mid;
+        }
+        high--;
+        low = mid + 1;
+    }
+    else { /* cp < aray[mid] */
+        if (cp < array[0]) { /* Fail if outside the array */
+            return -1;
+        }
+        high = mid;
+        if (cp >= array[mid - 1]) {
+            goto found_entry;
+        }
+    }
+
      /* Binary search.  What we are looking for is <i> such that
       * array[i] <= cp < array[i+1]
-     * The loop below converges on the i+1. */
+     * The loop below converges on the i+1.  Note that there may not be an
+     * (i+1)th element in the array, and things work nonetheless */
      while (low < high) {
-       IV mid = (low + high) / 2;
-       if (array[mid] <= cp) {
+       mid = (low + high) / 2;
+        assert(mid <= highest_element);
+       if (array[mid] <= cp) { /* cp >= array[mid] */
             low = mid + 1;
  
             /* We could do this extra test to exit the loop early.
@@ -7356,7 +7381,10 @@ Perl__invlist_search(pTHX_ SV* const invlist, const UV cp)
         }
      }
  
-    return high - 1;
+  found_entry:
+    high--;
+    invlist_set_previous_index(invlist, high);
+    return high;
  }
  
  void
@@ -7370,7 +7398,7 @@ Perl__invlist_populate_swatch(pTHX_ SV* const invlist, const UV start, const UV
       * that <swatch> is all 0's on input */
  
      UV current = start;
-    const IV len = invlist_len(invlist);
+    const IV len = _invlist_len(invlist);
      IV i;
      const UV * array;
  
@@ -7402,7 +7430,15 @@ Perl__invlist_populate_swatch(pTHX_ SV* const invlist, const UV start, const UV
              current = array[i];
             if (current >= end) {   /* Finished if beyond the end of what we
                                        are populating */
-                return;
+                if (LIKELY(end < UV_MAX)) {
+                    return;
+                }
+
+                /* We get here when the upper bound is the maximum
+                 * representable on the machine, and we are looking for just
+                 * that code point.  Have to special case it */
+                i = len;
+                goto join_end_of_list;
              }
          }
          assert(current >= start);
@@ -7419,6 +7455,8 @@ Perl__invlist_populate_swatch(pTHX_ SV* const invlist, const UV start, const UV
              swatch[offset >> 3] |= 1 << (offset & 7);
          }
  
+    join_end_of_list:
+
         /* Quit if at the end of the list */
          if (i >= len) {
  
@@ -7490,7 +7528,7 @@ Perl__invlist_union_maybe_complement_2nd(pTHX_ SV* const a, SV* const b, bool co
      assert(a != b);
  
      /* If either one is empty, the union is the other one */
-    if (a == NULL || ((len_a = invlist_len(a)) == 0)) {
+    if (a == NULL || ((len_a = _invlist_len(a)) == 0)) {
         if (*output == a) {
              if (a != NULL) {
                  SvREFCNT_dec(a);
@@ -7504,7 +7542,7 @@ Perl__invlist_union_maybe_complement_2nd(pTHX_ SV* const a, SV* const b, bool co
         } /* else *output already = b; */
         return;
      }
-    else if ((len_b = invlist_len(b)) == 0) {
+    else if ((len_b = _invlist_len(b)) == 0) {
         if (*output == b) {
             SvREFCNT_dec(b);
         }
@@ -7645,7 +7683,7 @@ Perl__invlist_union_maybe_complement_2nd(pTHX_ SV* const a, SV* const b, bool co
  
      /* Set result to final length, which can change the pointer to array_u, so
       * re-find it */
-    if (len_u != invlist_len(u)) {
+    if (len_u != _invlist_len(u)) {
         invlist_set_len(u, len_u);
         invlist_trim(u);
         array_u = invlist_array(u);
@@ -7724,8 +7762,8 @@ Perl__invlist_intersection_maybe_complement_2nd(pTHX_ SV* const a, SV* const b,
      assert(a != b);
  
      /* Special case if either one is empty */
-    len_a = invlist_len(a);
-    if ((len_a == 0) || ((len_b = invlist_len(b)) == 0)) {
+    len_a = _invlist_len(a);
+    if ((len_a == 0) || ((len_b = _invlist_len(b)) == 0)) {
  
          if (len_a != 0 && complement_b) {
  
@@ -7871,7 +7909,7 @@ Perl__invlist_intersection_maybe_complement_2nd(pTHX_ SV* const a, SV* const b,
  
      /* Set result to final length, which can change the pointer to array_r, so
       * re-find it */
-    if (len_r != invlist_len(r)) {
+    if (len_r != _invlist_len(r)) {
         invlist_set_len(r, len_r);
         invlist_trim(r);
         array_r = invlist_array(r);
@@ -7919,13 +7957,13 @@ Perl__add_range_to_invlist(pTHX_ SV* invlist, const UV start, const UV end)
         len = 0;
      }
      else {
-       len = invlist_len(invlist);
+       len = _invlist_len(invlist);
      }
  
      /* If comes after the final entry, can just append it to the end */
      if (len == 0
         || start >= invlist_array(invlist)
-                                   [invlist_len(invlist) - 1])
+                                   [_invlist_len(invlist) - 1])
      {
         _append_range_to_invlist(invlist, start, end);
         return invlist;
@@ -7946,18 +7984,6 @@ Perl__add_range_to_invlist(pTHX_ SV* invlist, const UV start, const UV end)
  
  #endif
  
-PERL_STATIC_INLINE bool
-S__invlist_contains_cp(pTHX_ SV* const invlist, const UV cp)
-{
-    /* Does <invlist> contain code point <cp> as part of the set? */
-
-    IV index = _invlist_search(invlist, cp);
-
-    PERL_ARGS_ASSERT__INVLIST_CONTAINS_CP;
-
-    return index >= 0 && ELEMENT_RANGE_MATCHES_INVLIST(index);
-}
-
  PERL_STATIC_INLINE SV*
  S_add_cp_to_invlist(pTHX_ SV* invlist, const UV cp) {
      return _add_range_to_invlist(invlist, cp, cp);
@@ -7971,7 +7997,7 @@ Perl__invlist_invert(pTHX_ SV* const invlist)
       * have a zero; removes it otherwise.  As described above, the data
       * structure is set up so that this is very efficient */
  
-    UV* len_pos = get_invlist_len_addr(invlist);
+    UV* len_pos = _get_invlist_len_addr(invlist);
  
      PERL_ARGS_ASSERT__INVLIST_INVERT;
  
@@ -8008,7 +8034,7 @@ Perl__invlist_invert_prop(pTHX_ SV* const invlist)
  
      _invlist_invert(invlist);
  
-    len = invlist_len(invlist);
+    len = _invlist_len(invlist);
  
      if (len != 0) { /* If empty do nothing */
         array = invlist_array(invlist);
@@ -8040,7 +8066,7 @@ S_invlist_clone(pTHX_ SV* const invlist)
  
      /* Need to allocate extra space to accommodate Perl's addition of a
       * trailing NUL to SvPV's, since it thinks they are always strings */
-    SV* new_invlist = _new_invlist(invlist_len(invlist) + 1);
+    SV* new_invlist = _new_invlist(_invlist_len(invlist) + 1);
      STRLEN length = SvCUR(invlist);
  
      PERL_ARGS_ASSERT_INVLIST_CLONE;
@@ -8091,7 +8117,7 @@ S_invlist_iternext(pTHX_ SV* invlist, UV* start, UV* end)
       * will start over at the beginning of the list */
  
      UV* pos = get_invlist_iter_addr(invlist);
-    UV len = invlist_len(invlist);
+    UV len = _invlist_len(invlist);
      UV *array;
  
      PERL_ARGS_ASSERT_INVLIST_ITERNEXT;
@@ -8123,7 +8149,7 @@ S_invlist_highest(pTHX_ SV* const invlist)
       * 0, or if the list is empty.  If this distinction matters to you, check
       * for emptiness before calling this function */
  
-    UV len = invlist_len(invlist);
+    UV len = _invlist_len(invlist);
      UV *array;
  
      PERL_ARGS_ASSERT_INVLIST_HIGHEST;
@@ -8210,8 +8236,8 @@ S__invlistEQ(pTHX_ SV* const a, SV* const b, bool complement_b)
  
      UV* array_a = invlist_array(a);
      UV* array_b = invlist_array(b);
-    UV len_a = invlist_len(a);
-    UV len_b = invlist_len(b);
+    UV len_a = _invlist_len(a);
+    UV len_b = _invlist_len(b);
  
      UV i = 0;              /* current index into the arrays */
      bool retval = TRUE;     /* Assume are identical until proven otherwise */
@@ -8303,11 +8329,11 @@ S_reg(pTHX_ RExC_state_t *pRExC_state, I32 paren, I32 *flagp,U32 depth)
      /* paren: Parenthesized? 0=top, 1=(, inside: changed to letter. */
  {
      dVAR;
-    register regnode *ret;             /* Will be the head of the group. */
-    register regnode *br;
-    register regnode *lastbr;
-    register regnode *ender = NULL;
-    register I32 parno = 0;
+    regnode *ret;              /* Will be the head of the group. */
+    regnode *br;
+    regnode *lastbr;
+    regnode *ender = NULL;
+    I32 parno = 0;
      I32 flags;
      U32 oregflags = RExC_flags;
      bool have_branch = 0;
@@ -9313,9 +9339,9 @@ STATIC regnode *
  S_regbranch(pTHX_ RExC_state_t *pRExC_state, I32 *flagp, I32 first, U32 depth)
  {
      dVAR;
-    register regnode *ret;
-    register regnode *chain = NULL;
-    register regnode *latest;
+    regnode *ret;
+    regnode *chain = NULL;
+    regnode *latest;
      I32 flags = 0, c = 0;
      GET_RE_DEBUG_FLAGS_DECL;
  
@@ -9386,9 +9412,9 @@ STATIC regnode *
  S_regpiece(pTHX_ RExC_state_t *pRExC_state, I32 *flagp, U32 depth)
  {
      dVAR;
-    register regnode *ret;
-    register char op;
-    register char *next;
+    regnode *ret;
+    char op;
+    char *next;
      I32 flags;
      const char * const origparse = RExC_parse;
      I32 min;
@@ -9577,20 +9603,22 @@ S_regpiece(pTHX_ RExC_state_t *pRExC_state, I32 *flagp, U32 depth)
      return(ret);
  }
  
-/* grok_bslash_N(pRExC_state, regnode** node_p, UV *valuep, UV depth, bool in_charclass)
+STATIC bool
+S_grok_bslash_N(pTHX_ RExC_state_t *pRExC_state, regnode** node_p, UV *valuep, I32 *flagp, U32 depth, bool in_char_class)
+{
     
-   This is expected to be called by a parser routine that has recognized '\N'
+ /* This is expected to be called by a parser routine that has recognized '\N'
     and needs to handle the rest. RExC_parse is expected to point at the first
     char following the N at the time of the call.  On successful return,
     RExC_parse has been updated to point to just after the sequence identified
-   by this routine.
+   by this routine, and <*flagp> has been updated.
  
-   The \N may be inside (indicated by the boolean <in_charclass>) or outside a
+   The \N may be inside (indicated by the boolean <in_char_class>) or outside a
     character class.
  
     \N may begin either a named sequence, or if outside a character class, mean
     to match a non-newline.  For non single-quoted regexes, the tokenizer has
-   attempted to decide which, and in the case of a named sequence converted it
+   attempted to decide which, and in the case of a named sequence, converted it
     into one of the forms: \N{} (if the sequence is null), or \N{U+c1.c2...},
     where c1... are the characters in the sequence.  For single-quoted regexes,
     the tokenizer passes the \N sequence through unchanged; this code will not
@@ -9622,9 +9650,6 @@ S_regpiece(pTHX_ RExC_state_t *pRExC_state, I32 *flagp, U32 depth)
     null.
   */
  
-STATIC bool
-S_grok_bslash_N(pTHX_ RExC_state_t *pRExC_state, regnode** node_p, UV *valuep, I32 *flagp, U32 depth, bool in_char_class)
-{
      char * endbrace;    /* '}' following the name */
      char* p;
      char *endchar;     /* Points to '.' or '}' ending cur char in the input
@@ -9774,6 +9799,7 @@ S_grok_bslash_N(pTHX_ RExC_state_t *pRExC_state, regnode** node_p, UV *valuep, I
         SV * substitute_parse = newSVpvn_flags("?:", 2, SVf_UTF8|SVs_TEMP);
         STRLEN len;
         char *orig_end = RExC_end;
+        I32 flags;
  
         while (RExC_parse < endbrace) {
  
@@ -9799,7 +9825,8 @@ S_grok_bslash_N(pTHX_ RExC_state_t *pRExC_state, regnode** node_p, UV *valuep, I
         /* The values are Unicode, and therefore not subject to recoding */
         RExC_override_recoding = 1;
  
-       *node_p = reg(pRExC_state, 1, flagp, depth+1);
+       *node_p = reg(pRExC_state, 1, &flags, depth+1);
+       *flagp |= flags&(HASWIDTH|SPSTART|SIMPLE|POSTPONED);
  
         RExC_parse = endbrace;
         RExC_end = orig_end;
@@ -9866,15 +9893,23 @@ S_compute_EXACTish(pTHX_ RExC_state_t *pRExC_state)
  }
  
  PERL_STATIC_INLINE void
-S_alloc_maybe_populate_EXACT(pTHX_ RExC_state_t *pRExC_state, regnode *node, STRLEN len, UV code_point)
+S_alloc_maybe_populate_EXACT(pTHX_ RExC_state_t *pRExC_state, regnode *node, I32* flagp, STRLEN len, UV code_point)
  {
-    /* This knows the details about sizing an EXACTish node, and potentially
-     * populating it with a single character.  If <len> is non-zero, it assumes
-     * that the node has already been populated, and just does the sizing,
-     * ignoring <code_point>.  Otherwise it looks at <code_point> and
-     * calculates what <len> should be.  In pass 1, it sizes the node
-     * appropriately.  In pass 2, it additionally will populate the node's
-     * STRING with <code_point>, if <len> is 0.
+    /* This knows the details about sizing an EXACTish node, setting flags for
+     * it (by setting <*flagp>, and potentially populating it with a single
+     * character.
+     *
+     * If <len> is non-zero, this function assumes that the node has already
+     * been populated, and just does the sizing.  In this case <code_point>
+     * should be the final code point that has already been placed into the
+     * node.  This value will be ignored except that under some circumstances
+     * <*flagp> is set based on it.
+     *
+     * If <len is zero, the function assumes that the node is to contain only
+     * the single character given by <code_point> and calculates what <len>
+     * should be.  In pass 1, it sizes the node appropriately.  In pass 2, it
+     * additionally will populate the node's STRING with <code_point>, if <len>
+     * is 0.  In both cases <*flagp> is appropriately set
       *
       * It knows that under FOLD, UTF characters and the Latin Sharp S must be
       * folded (the latter only when the rules indicate it can match 'ss') */
@@ -9919,6 +9954,10 @@ S_alloc_maybe_populate_EXACT(pTHX_ RExC_state_t *pRExC_state, regnode *node, STR
              Copy((char *) character, STRING(node), len, char);
          }
      }
+
+    *flagp |= HASWIDTH;
+    if (len == 1 && UNI_IS_INVARIANT(code_point))
+        *flagp |= SIMPLE;
  }
  
  /*
@@ -10033,13 +10072,12 @@ tryagain:
      case '[':
      {
         char * const oregcomp_parse = ++RExC_parse;
-        ret = regclass(pRExC_state,depth+1);
+        ret = regclass(pRExC_state, flagp,depth+1);
         if (*RExC_parse != ']') {
             RExC_parse = oregcomp_parse;
             vFAIL("Unmatched [");
         }
         nextchar(pRExC_state);
-       *flagp |= HASWIDTH|SIMPLE;
          Set_Node_Length(ret, RExC_parse - oregcomp_parse + 1); /* MJD */
         break;
      }
@@ -10250,7 +10288,7 @@ tryagain:
                 }
                 RExC_parse--;
  
-                ret = regclass(pRExC_state,depth+1);
+                ret = regclass(pRExC_state, flagp,depth+1);
  
                 RExC_end = oldregxend;
                 RExC_parse--;
@@ -10258,7 +10296,6 @@ tryagain:
                 Set_Node_Offset(ret, parse_start + 2);
                 Set_Node_Cur_Length(ret);
                 nextchar(pRExC_state);
-               *flagp |= HASWIDTH|SIMPLE;
             }
             break;
          case 'N': 
@@ -10421,9 +10458,9 @@ tryagain:
             RExC_parse++;
  
         defchar: {
-           register STRLEN len = 0;
+           STRLEN len = 0;
             UV ender;
-           register char *p;
+           char *p;
             char *s;
  #define MAX_NODE_STRING_SIZE 127
             char foldbuf[MAX_NODE_STRING_SIZE+UTF8_MAXBYTES_CASE];
@@ -10432,7 +10469,7 @@ tryagain:
             STRLEN foldlen;
              U8 node_type;
              bool next_is_quantifier;
-            char * oldp;
+            char * oldp = NULL;
  
             ender = 0;
              node_type = compute_EXACTish(pRExC_state);
@@ -10935,6 +10972,17 @@ tryagain:
  
         loopdone:   /* Jumped to when encounters something that shouldn't be in
                        the node */
+
+            /* I (khw) don't know if you can get here with zero length, but the
+             * old code handled this situation by creating a zero-length EXACT
+             * node.  Might as well be NOTHING instead */
+            if (len == 0) {
+                OP(ret) = NOTHING;
+            }
+            else{
+                alloc_maybe_populate_EXACT(pRExC_state, ret, flagp, len, ender);
+            }
+
             RExC_parse = p - 1;
              Set_Node_Cur_Length(ret); /* MJD */
             nextchar(pRExC_state);
@@ -10944,12 +10992,7 @@ tryagain:
                 if (iv < 0)
                     vFAIL("Internal disaster");
             }
-           if (len > 0)
-               *flagp |= HASWIDTH;
-           if (len == 1 && UNI_IS_INVARIANT(ender))
-               *flagp |= SIMPLE;
  
-            alloc_maybe_populate_EXACT(pRExC_state, ret, len, 0);
         } /* End of label 'defchar:' */
         break;
      } /* End of giant switch on input character */
@@ -11316,14 +11359,14 @@ S_add_alternate(pTHX_ AV** alternate_ptr, U8* string, STRLEN len)
     above 255, a range list is used */
  
  STATIC regnode *
-S_regclass(pTHX_ RExC_state_t *pRExC_state, U32 depth)
+S_regclass(pTHX_ RExC_state_t *pRExC_state, I32 *flagp, U32 depth)
  {
      dVAR;
-    register UV nextvalue;
-    register UV prevvalue = OOB_UNICODE;
-    register IV range = 0;
+    UV nextvalue;
+    UV prevvalue = OOB_UNICODE;
+    IV range = 0;
      UV value = 0;
-    register regnode *ret;
+    regnode *ret;
      STRLEN numlen;
      IV namedclass = OOB_NAMEDCLASS;
      char *rangebegin = NULL;
@@ -11487,7 +11530,7 @@ parseit:
                      if this makes sense as it does change the behaviour
                      from earlier versions, OTOH that behaviour was broken
                      as well. */
-                    if (! grok_bslash_N(pRExC_state, NULL, &value, NULL, depth,
+                    if (! grok_bslash_N(pRExC_state, NULL, &value, flagp, depth,
                                        TRUE /* => charclass */))
                      {
                          goto parseit;
@@ -11993,8 +12036,8 @@ parseit:
                     }
                      if (!SIZE_ONLY) {
                          cp_list = add_cp_to_invlist(cp_list, '-');
-                        element_count++;
                      }
+                    element_count++;
                 } else
                     range = 1;  /* yeah, it's a range! */
                 continue;       /* but do it the next time */
@@ -12106,6 +12149,7 @@ parseit:
                      if (invert) {
                          op += NALNUM - ALNUM;
                      }
+                    *flagp |= HASWIDTH|SIMPLE;
                      break;
  
                  /* The second group doesn't depend of the charset modifiers.
@@ -12116,6 +12160,7 @@ parseit:
                  case ANYOF_HORIZWS:
                    is_horizws:
                      op = (invert) ? NHORIZWS : HORIZWS;
+                    *flagp |= HASWIDTH|SIMPLE;
                      break;
  
                  case ANYOF_NVERTWS:
@@ -12123,6 +12168,7 @@ parseit:
                      /* FALLTHROUGH */
                  case ANYOF_VERTWS:
                      op = (invert) ? NVERTWS : VERTWS;
+                    *flagp |= HASWIDTH|SIMPLE;
                      break;
  
                  case ANYOF_MAX:
@@ -12162,6 +12208,8 @@ parseit:
              if (invert) {
                  if (! LOC && value == '\n') {
                      op = REG_ANY; /* Optimize [^\n] */
+                    *flagp |= HASWIDTH|SIMPLE;
+                    RExC_naughty++;
                  }
              }
              else if (value < 256 || UTF) {
@@ -12175,6 +12223,7 @@ parseit:
              if (prevvalue == '0') {
                  if (value == '9') {
                      op = (invert) ? NDIGITA : DIGITA;
+                    *flagp |= HASWIDTH|SIMPLE;
                  }
              }
          }
@@ -12208,9 +12257,10 @@ parseit:
                  if (! SIZE_ONLY) {
                      FLAGS(ret) = arg;
                  }
+                *flagp |= HASWIDTH|SIMPLE;
              }
              else if (PL_regkind[op] == EXACT) {
-                alloc_maybe_populate_EXACT(pRExC_state, ret, 0, value);
+                alloc_maybe_populate_EXACT(pRExC_state, ret, flagp, 0, value);
              }
  
              RExC_parse = (char *) cur_parse;
@@ -12222,7 +12272,7 @@ parseit:
  
      if (SIZE_ONLY)
          return ret;
-    /****** !SIZE_ONLY AFTER HERE *********/
+    /****** !SIZE_ONLY (Pass 2) AFTER HERE *********/
  
      /* If folding, we calculate all characters that could fold to or from the
       * ones already on the list */
@@ -12260,7 +12310,7 @@ parseit:
                   * rules hard-coded into Perl.  (This case happens legitimately
                   * during compilation of Perl itself before the Unicode tables
                   * are generated) */
-                if (invlist_len(PL_utf8_foldable) == 0) {
+                if (_invlist_len(PL_utf8_foldable) == 0) {
                      PL_utf8_foldclosures = newHV();
                  }
                  else {
@@ -12678,6 +12728,7 @@ parseit:
               * it doesn't match anything.  (perluniprops.pod notes such
               * properties) */
              op = OPFAIL;
+            *flagp |= HASWIDTH|SIMPLE;
          }
          else if (start == end) {    /* The range is a single code point */
              if (! invlist_iternext(cp_list, &start, &end)
@@ -12743,12 +12794,16 @@ parseit:
          else if (start == 0) {
              if (end == UV_MAX) {
                  op = SANY;
+                *flagp |= HASWIDTH|SIMPLE;
+                RExC_naughty++;
              }
              else if (end == '\n' - 1
                      && invlist_iternext(cp_list, &start, &end)
                      && start == '\n' + 1 && end == UV_MAX)
              {
                  op = REG_ANY;
+                *flagp |= HASWIDTH|SIMPLE;
+                RExC_naughty++;
              }
          }
  
@@ -12761,7 +12816,7 @@ parseit:
              RExC_parse = (char *)cur_parse;
  
              if (PL_regkind[op] == EXACT) {
-                alloc_maybe_populate_EXACT(pRExC_state, ret, 0, value);
+                alloc_maybe_populate_EXACT(pRExC_state, ret, flagp, 0, value);
              }
  
              SvREFCNT_dec(listsv);
@@ -12817,7 +12872,7 @@ parseit:
         }
  
         /* If have completely emptied it, remove it completely */
-       if (invlist_len(cp_list) == 0) {
+       if (_invlist_len(cp_list) == 0) {
             SvREFCNT_dec(cp_list);
             cp_list = NULL;
         }
@@ -12902,6 +12957,8 @@ parseit:
         RExC_rxi->data->data[n] = (void*)rv;
         ARG_SET(ret, n);
      }
+
+    *flagp |= HASWIDTH|SIMPLE;
      return ret;
  }
  #undef HAS_NONLOCALE_RUNTIME_PROPERTY_DEFINITION
@@ -12995,7 +13052,7 @@ STATIC regnode *                        /* Location. */
  S_reg_node(pTHX_ RExC_state_t *pRExC_state, U8 op)
  {
      dVAR;
-    register regnode *ptr;
+    regnode *ptr;
      regnode * const ret = RExC_emit;
      GET_RE_DEBUG_FLAGS_DECL;
  
@@ -13037,7 +13094,7 @@ STATIC regnode *                        /* Location. */
  S_reganode(pTHX_ RExC_state_t *pRExC_state, U8 op, U32 arg)
  {
      dVAR;
-    register regnode *ptr;
+    regnode *ptr;
      regnode * const ret = RExC_emit;
      GET_RE_DEBUG_FLAGS_DECL;
  
@@ -13109,9 +13166,9 @@ STATIC void
  S_reginsert(pTHX_ RExC_state_t *pRExC_state, U8 op, regnode *opnd, U32 depth)
  {
      dVAR;
-    register regnode *src;
-    register regnode *dst;
-    register regnode *place;
+    regnode *src;
+    regnode *dst;
+    regnode *place;
      const int offset = regarglen[(U8)op];
      const int size = NODE_STEP_REGNODE + offset;
      GET_RE_DEBUG_FLAGS_DECL;
@@ -13197,7 +13254,7 @@ STATIC void
  S_regtail(pTHX_ RExC_state_t *pRExC_state, regnode *p, const regnode *val,U32 depth)
  {
      dVAR;
-    register regnode *scan;
+    regnode *scan;
      GET_RE_DEBUG_FLAGS_DECL;
  
      PERL_ARGS_ASSERT_REGTAIL;
@@ -13256,7 +13313,7 @@ STATIC U8
  S_regtail_study(pTHX_ RExC_state_t *pRExC_state, regnode *p, const regnode *val,U32 depth)
  {
      dVAR;
-    register regnode *scan;
+    regnode *scan;
      U8 exact = PSEUDO;
  #ifdef EXPERIMENTAL_INPLACESCAN
      I32 min = 0;
@@ -13497,7 +13554,7 @@ Perl_regprop(pTHX_ const regexp *prog, SV *sv, const regnode *o)
  {
  #ifdef DEBUGGING
      dVAR;
-    register int k;
+    int k;
  
      /* Should be synchronized with * ANYOF_ #xdefines in regcomp.h */
      static const char * const anyofs[] = {
@@ -14301,7 +14358,7 @@ regnode *
  Perl_regnext(pTHX_ register regnode *p)
  {
      dVAR;
-    register I32 offset;
+    I32 offset;
  
      if (!p)
         return(NULL);
@@ -14460,8 +14517,8 @@ S_dumpuntil(pTHX_ const regexp *r, const regnode *start, const regnode *node,
             SV* sv, I32 indent, U32 depth)
  {
      dVAR;
-    register U8 op = PSEUDO;   /* Arbitrary non-END op. */
-    register const regnode *next;
+    U8 op = PSEUDO;    /* Arbitrary non-END op. */
+    const regnode *next;
      const regnode *optstart= NULL;
      
      RXi_GET_DECL(r,ri);
@@ -14512,9 +14569,9 @@ S_dumpuntil(pTHX_ const regexp *r, const regnode *start, const regnode *node,
         if (PL_regkind[(U8)op] == BRANCHJ) {
             assert(next);
             {
-                register const regnode *nnode = (OP(next) == LONGJMP
-                                            ? regnext((regnode *)next)
-                                            : next);
+                const regnode *nnode = (OP(next) == LONGJMP
+                                       ? regnext((regnode *)next)
+                                       : next);
                  if (last && nnode > last)
                      nnode = last;
                  DUMPUNTIL(NEXTOPER(NEXTOPER(node)), nnode);