-/* $Id: extract.c,v 1.267 2007-10-31 16:56:14 adam Exp $
+/* $Id: extract.c,v 1.270 2007-11-30 12:19:08 adam Exp $
Copyright (C) 1995-2007
Index Data ApS
int cmd, zebra_rec_keys_t skp);
static void extract_schema_add(struct recExtractCtrl *p, Odr_oid *oid);
static void extract_token_add(RecWord *p);
-static void extract_token_add2(RecWord *p);
static void check_log_limit(ZebraHandle zh)
{
struct recordGroup *rGroup;
};
-static void all_matches_add(struct recExtractCtrl *ctrl)
+/** \brief add the always-matches index entry and map to real record ID
+ \param ctrl record control
+ \param record_id custom record ID
+ \param sysno system record ID
+
+ This function serves two purposes.. It adds the always matches
+ entry and makes a pointer from the custom record ID (if defined)
+ back to the system record ID (sysno)
+ See zebra_recid_to_sysno .
+ */
+static void all_matches_add(struct recExtractCtrl *ctrl, zint record_id,
+ zint sysno)
{
RecWord word;
extract_init(ctrl, &word);
+ word.record_id = record_id;
+ /* we use the seqno as placeholder for a way to get back to
+ record database from _ALLRECORDS.. This is used if a custom
+ RECORD was defined */
+ word.seqno = sysno;
word.index_name = "_ALLRECORDS";
word.index_type = "w";
- word.seqno = 1;
+
extract_add_index_string(&word, zinfo_index_category_alwaysmatches,
"", 0);
}
stream->endf(stream, &null_offset);;
extractCtrl.init = extract_init;
- if (zh->reg->index_types)
- {
- extractCtrl.tokenAdd = extract_token_add2;
- }
- else
- {
- extractCtrl.tokenAdd = extract_token_add;
- }
+ extractCtrl.tokenAdd = extract_token_add;
extractCtrl.schemaAdd = extract_schema_add;
extractCtrl.dh = zh->reg->dh;
extractCtrl.handle = zh;
else
end_offset = stream->tellf(stream);
- all_matches_add(&extractCtrl);
-
if (extractCtrl.match_criteria[0])
match_criteria = extractCtrl.match_criteria;
}
}
}
}
+
if (zebra_rec_keys_empty(zh->reg->keys))
{
/* the extraction process returned no information - the record
*sysno = rec->sysno;
+
+ if (stream)
+ {
+ all_matches_add(&extractCtrl,
+ zebra_rec_keys_get_custom_record_id(zh->reg->keys),
+ *sysno);
+ }
+
+
recordAttr = rec_init_attr(zh->reg->zei, rec);
if (extractCtrl.staticrank < 0)
{
rec = rec_get(zh->reg->records, *sysno);
assert(rec);
+
+ if (stream)
+ {
+ all_matches_add(&extractCtrl,
+ zebra_rec_keys_get_custom_record_id(zh->reg->keys),
+ *sysno);
+ }
recordAttr = rec_init_attr(zh->reg->zei, rec);
extract_add_string(p, zm, buf, i);
}
-static void extract_token_add2_index(ZebraHandle zh, zebra_index_type_t type,
- RecWord *p)
+static void extract_add_icu(RecWord *p, zebra_map_t zm)
{
struct it_key key;
const char *res_buf = 0;
size_t res_len = 0;
- int r = zebra_index_type_tokenize(type, p->term_buf, p->term_len,
- &res_buf, &res_len);
+ ZebraHandle zh = p->extractCtrl->handle;
+ int r = zebra_map_tokenize(zm, p->term_buf, p->term_len,
+ &res_buf, &res_len);
int cat = zinfo_index_category_index;
int ch = zebraExplain_lookup_attr_str(zh->reg->zei, cat, p->index_type, p->index_name);
if (ch < 0)
key.mem[i++] = p->seqno;
key.len = i;
- yaz_log(YLOG_LOG, "keys_write %.*s", (int) res_len, res_buf);
zebra_rec_keys_write(zh->reg->keys, res_buf, res_len, &key);
p->seqno++;
- r = zebra_index_type_tokenize(type, 0, 0, &res_buf, &res_len);
+ r = zebra_map_tokenize(zm, 0, 0, &res_buf, &res_len);
}
}
-static void extract_token_add2(RecWord *p)
-{
- ZebraHandle zh = p->extractCtrl->handle;
- zebra_index_type_t type = zebra_index_type_get(zh->reg->index_types, p->index_type);
- if (type)
- {
- if (zebra_index_type_is_index(type))
- {
- extract_token_add2_index(zh, type, p);
- }
- else if (zebra_index_type_is_sort(type))
- {
- ;
-
- }
- }
-}
/** \brief top-level indexing handler for recctrl system
\param p token data to be indexed
}
if ((wrbuf = zebra_replace(zm, 0, p->term_buf, p->term_len)))
{
- p->term_buf = wrbuf_buf(wrbuf);
- p->term_len = wrbuf_len(wrbuf);
+ p->term_buf = wrbuf_buf(wrbuf);
+ p->term_len = wrbuf_len(wrbuf);
+ }
+ if (zebra_maps_is_icu(zm))
+ {
+ extract_add_icu(p, zm);
}
- if (zebra_maps_is_complete(zm))
- extract_add_complete_field(p, zm);
else
- extract_add_incomplete_field(p, zm);
+ {
+ if (zebra_maps_is_complete(zm))
+ extract_add_complete_field(p, zm);
+ else
+ extract_add_incomplete_field(p, zm);
+ }
}
static void extract_set_store_data_cb(struct recExtractCtrl *p,