2014-02-10 02:10:30 +01:00
|
|
|
/*************************************************************************/
|
|
|
|
/* compressed_translation.cpp */
|
|
|
|
/*************************************************************************/
|
|
|
|
/* This file is part of: */
|
|
|
|
/* GODOT ENGINE */
|
2017-08-27 14:16:55 +02:00
|
|
|
/* https://godotengine.org */
|
2014-02-10 02:10:30 +01:00
|
|
|
/*************************************************************************/
|
2020-01-01 11:16:22 +01:00
|
|
|
/* Copyright (c) 2007-2020 Juan Linietsky, Ariel Manzur. */
|
|
|
|
/* Copyright (c) 2014-2020 Godot Engine contributors (cf. AUTHORS.md). */
|
2014-02-10 02:10:30 +01:00
|
|
|
/* */
|
|
|
|
/* Permission is hereby granted, free of charge, to any person obtaining */
|
|
|
|
/* a copy of this software and associated documentation files (the */
|
|
|
|
/* "Software"), to deal in the Software without restriction, including */
|
|
|
|
/* without limitation the rights to use, copy, modify, merge, publish, */
|
|
|
|
/* distribute, sublicense, and/or sell copies of the Software, and to */
|
|
|
|
/* permit persons to whom the Software is furnished to do so, subject to */
|
|
|
|
/* the following conditions: */
|
|
|
|
/* */
|
|
|
|
/* The above copyright notice and this permission notice shall be */
|
|
|
|
/* included in all copies or substantial portions of the Software. */
|
|
|
|
/* */
|
|
|
|
/* THE SOFTWARE IS PROVIDED "AS IS", WITHOUT WARRANTY OF ANY KIND, */
|
|
|
|
/* EXPRESS OR IMPLIED, INCLUDING BUT NOT LIMITED TO THE WARRANTIES OF */
|
|
|
|
/* MERCHANTABILITY, FITNESS FOR A PARTICULAR PURPOSE AND NONINFRINGEMENT.*/
|
|
|
|
/* IN NO EVENT SHALL THE AUTHORS OR COPYRIGHT HOLDERS BE LIABLE FOR ANY */
|
|
|
|
/* CLAIM, DAMAGES OR OTHER LIABILITY, WHETHER IN AN ACTION OF CONTRACT, */
|
|
|
|
/* TORT OR OTHERWISE, ARISING FROM, OUT OF OR IN CONNECTION WITH THE */
|
|
|
|
/* SOFTWARE OR THE USE OR OTHER DEALINGS IN THE SOFTWARE. */
|
|
|
|
/*************************************************************************/
|
2018-01-05 00:50:27 +01:00
|
|
|
|
2014-02-10 02:10:30 +01:00
|
|
|
#include "compressed_translation.h"
|
2017-01-16 08:04:19 +01:00
|
|
|
|
2018-09-11 18:13:45 +02:00
|
|
|
#include "core/pair.h"
|
2014-02-10 02:10:30 +01:00
|
|
|
|
Split thirdparty smaz.c out of compressed_translation.cpp
Code comes from https://github.com/antirez/smaz/blob/150e125cbae2e8fd20dd332432776ce13395d4d4/smaz.c
With a small modification to match Godot expectations:
```
diff --git a/thirdparty/core/smaz.c b/thirdparty/core/smaz.c
index 9b1ebc2..555dfea 100644
--- a/thirdparty/core/smaz.c
+++ b/thirdparty/core/smaz.c
@@ -14,7 +14,7 @@ THIS SOFTWARE IS PROVIDED BY THE COPYRIGHT HOLDERS AND CONTRIBUTORS "AS IS" AND
#include <string.h>
/* Our compression codebook, used for compression */
-static char *Smaz_cb[241] = {
+static const char *Smaz_cb[241] = {
"\002s,\266", "\003had\232\002leW", "\003on \216", "", "\001yS",
"\002ma\255\002li\227", "\003or \260", "", "\002ll\230\003s t\277",
"\004fromg\002mel", "", "\003its\332", "\001z\333", "\003ingF", "\001>\336",
@@ -89,7 +89,7 @@ static char *Smaz_rcb[254] = {
"e, ", " it", "whi", " ma", "ge", "x", "e c", "men", ".com"
};
-int smaz_compress(char *in, int inlen, char *out, int outlen) {
+int smaz_compress(const char *in, int inlen, char *out, int outlen) {
unsigned int h1,h2,h3=0;
int verblen = 0, _outlen = outlen;
char verb[256], *_out = out;
@@ -167,7 +167,7 @@ out:
return out-_out;
}
-int smaz_decompress(char *in, int inlen, char *out, int outlen) {
+int smaz_decompress(const char *in, int inlen, char *out, int outlen) {
unsigned char *c = (unsigned char*) in;
char *_out = out;
int _outlen = outlen;
@@ -192,7 +192,7 @@ int smaz_decompress(char *in, int inlen, char *out, int outlen) {
inlen -= 2+len;
} else {
/* Codebook entry */
- char *s = Smaz_rcb[*c];
+ const char *s = Smaz_rcb[*c];
int len = strlen(s);
if (outlen < len) return _outlen+1;
diff --git a/thirdparty/core/smaz.h b/thirdparty/core/smaz.h
index a547d89..a9d8a33 100644
--- a/thirdparty/core/smaz.h
+++ b/thirdparty/core/smaz.h
@@ -14,7 +14,7 @@ THIS SOFTWARE IS PROVIDED BY THE COPYRIGHT HOLDERS AND CONTRIBUTORS "AS IS" AND
#ifndef _SMAZ_H
#define _SMAZ_H
-int smaz_compress(char *in, int inlen, char *out, int outlen);
-int smaz_decompress(char *in, int inlen, char *out, int outlen);
+int smaz_compress(const char *in, int inlen, char *out, int outlen);
+int smaz_decompress(const char *in, int inlen, char *out, int outlen);
#endif
```
2017-04-28 19:00:11 +02:00
|
|
|
extern "C" {
|
|
|
|
#include "thirdparty/misc/smaz.h"
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
|
|
|
|
struct _PHashTranslationCmp {
|
|
|
|
|
|
|
|
int orig_len;
|
|
|
|
CharString compressed;
|
|
|
|
int offset;
|
|
|
|
};
|
|
|
|
|
|
|
|
void PHashTranslation::generate(const Ref<Translation> &p_from) {
|
|
|
|
#ifdef TOOLS_ENABLED
|
|
|
|
List<StringName> keys;
|
|
|
|
p_from->get_message_list(&keys);
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
int size = Math::larger_prime(keys.size());
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2020-03-17 07:33:00 +01:00
|
|
|
Vector<Vector<Pair<int, CharString>>> buckets;
|
|
|
|
Vector<Map<uint32_t, int>> table;
|
2017-03-05 16:44:50 +01:00
|
|
|
Vector<uint32_t> hfunc_table;
|
|
|
|
Vector<_PHashTranslationCmp> compressed;
|
2014-02-10 02:10:30 +01:00
|
|
|
|
|
|
|
table.resize(size);
|
|
|
|
hfunc_table.resize(size);
|
|
|
|
buckets.resize(size);
|
|
|
|
compressed.resize(keys.size());
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
int idx = 0;
|
|
|
|
int total_compression_size = 0;
|
|
|
|
int total_string_size = 0;
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
for (List<StringName>::Element *E = keys.front(); E; E = E->next()) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
|
|
|
//hash string
|
|
|
|
CharString cs = E->get().operator String().utf8();
|
2017-03-05 16:44:50 +01:00
|
|
|
uint32_t h = hash(0, cs.get_data());
|
|
|
|
Pair<int, CharString> p;
|
|
|
|
p.first = idx;
|
|
|
|
p.second = cs;
|
2018-07-25 03:11:03 +02:00
|
|
|
buckets.write[h % size].push_back(p);
|
2014-02-10 02:10:30 +01:00
|
|
|
|
|
|
|
//compress string
|
|
|
|
CharString src_s = p_from->get_message(E->get()).operator String().utf8();
|
|
|
|
_PHashTranslationCmp ps;
|
2017-03-05 16:44:50 +01:00
|
|
|
ps.orig_len = src_s.size();
|
|
|
|
ps.offset = total_compression_size;
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
if (ps.orig_len != 0) {
|
2014-02-10 02:10:30 +01:00
|
|
|
CharString dst_s;
|
|
|
|
dst_s.resize(src_s.size());
|
2018-12-16 01:44:18 +01:00
|
|
|
int ret = smaz_compress(src_s.get_data(), src_s.size(), dst_s.ptrw(), src_s.size());
|
2017-03-05 16:44:50 +01:00
|
|
|
if (ret >= src_s.size()) {
|
2014-02-10 02:10:30 +01:00
|
|
|
//if compressed is larger than original, just use original
|
2017-03-05 16:44:50 +01:00
|
|
|
ps.orig_len = src_s.size();
|
|
|
|
ps.compressed = src_s;
|
2014-02-10 02:10:30 +01:00
|
|
|
} else {
|
|
|
|
dst_s.resize(ret);
|
|
|
|
//ps.orig_len=;
|
2017-03-05 16:44:50 +01:00
|
|
|
ps.compressed = dst_s;
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
} else {
|
2017-03-05 16:44:50 +01:00
|
|
|
ps.orig_len = 1;
|
2014-02-10 02:10:30 +01:00
|
|
|
ps.compressed.resize(1);
|
2017-03-05 16:44:50 +01:00
|
|
|
ps.compressed[0] = 0;
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
|
2018-07-25 03:11:03 +02:00
|
|
|
compressed.write[idx] = ps;
|
2017-03-05 16:44:50 +01:00
|
|
|
total_compression_size += ps.compressed.size();
|
|
|
|
total_string_size += src_s.size();
|
2014-02-10 02:10:30 +01:00
|
|
|
idx++;
|
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
int bucket_table_size = 0;
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
for (int i = 0; i < size; i++) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2020-03-17 07:33:00 +01:00
|
|
|
const Vector<Pair<int, CharString>> &b = buckets[i];
|
2018-07-25 03:11:03 +02:00
|
|
|
Map<uint32_t, int> &t = table.write[i];
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
if (b.size() == 0)
|
2014-02-10 02:10:30 +01:00
|
|
|
continue;
|
|
|
|
|
|
|
|
int d = 1;
|
2017-03-05 16:44:50 +01:00
|
|
|
int item = 0;
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
while (item < b.size()) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
uint32_t slot = hash(d, b[item].second.get_data());
|
2014-02-10 02:10:30 +01:00
|
|
|
if (t.has(slot)) {
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
item = 0;
|
2014-02-10 02:10:30 +01:00
|
|
|
d++;
|
|
|
|
t.clear();
|
|
|
|
} else {
|
2017-03-05 16:44:50 +01:00
|
|
|
t[slot] = b[item].first;
|
2014-02-10 02:10:30 +01:00
|
|
|
item++;
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2018-07-25 03:11:03 +02:00
|
|
|
hfunc_table.write[i] = d;
|
2017-03-05 16:44:50 +01:00
|
|
|
bucket_table_size += 2 + b.size() * 4;
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
|
2019-10-14 11:40:55 +02:00
|
|
|
ERR_FAIL_COND(bucket_table_size == 0);
|
|
|
|
|
2014-02-10 02:10:30 +01:00
|
|
|
hash_table.resize(size);
|
|
|
|
bucket_table.resize(bucket_table_size);
|
|
|
|
|
2020-02-17 22:06:54 +01:00
|
|
|
int *htwb = hash_table.ptrw();
|
|
|
|
int *btwb = bucket_table.ptrw();
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
uint32_t *htw = (uint32_t *)&htwb[0];
|
|
|
|
uint32_t *btw = (uint32_t *)&btwb[0];
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
int btindex = 0;
|
|
|
|
int collisions = 0;
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
for (int i = 0; i < size; i++) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2018-07-25 03:11:03 +02:00
|
|
|
const Map<uint32_t, int> &t = table[i];
|
2017-03-05 16:44:50 +01:00
|
|
|
if (t.size() == 0) {
|
|
|
|
htw[i] = 0xFFFFFFFF; //nothing
|
2014-02-10 02:10:30 +01:00
|
|
|
continue;
|
2017-03-05 16:44:50 +01:00
|
|
|
} else if (t.size() > 1) {
|
|
|
|
collisions += t.size() - 1;
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
htw[i] = btindex;
|
|
|
|
btw[btindex++] = t.size();
|
|
|
|
btw[btindex++] = hfunc_table[i];
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
for (Map<uint32_t, int>::Element *E = t.front(); E; E = E->next()) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
btw[btindex++] = E->key();
|
|
|
|
btw[btindex++] = compressed[E->get()].offset;
|
|
|
|
btw[btindex++] = compressed[E->get()].compressed.size();
|
|
|
|
btw[btindex++] = compressed[E->get()].orig_len;
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
}
|
|
|
|
|
|
|
|
strings.resize(total_compression_size);
|
2020-02-17 22:06:54 +01:00
|
|
|
uint8_t *cw = strings.ptrw();
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
for (int i = 0; i < compressed.size(); i++) {
|
|
|
|
memcpy(&cw[compressed[i].offset], compressed[i].compressed.get_data(), compressed[i].compressed.size());
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
ERR_FAIL_COND(btindex != bucket_table_size);
|
2014-02-10 02:10:30 +01:00
|
|
|
set_locale(p_from->get_locale());
|
|
|
|
|
|
|
|
#endif
|
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
bool PHashTranslation::_set(const StringName &p_name, const Variant &p_value) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
|
|
|
String name = p_name.operator String();
|
2017-03-05 16:44:50 +01:00
|
|
|
if (name == "hash_table") {
|
|
|
|
hash_table = p_value;
|
|
|
|
} else if (name == "bucket_table") {
|
|
|
|
bucket_table = p_value;
|
|
|
|
} else if (name == "strings") {
|
|
|
|
strings = p_value;
|
|
|
|
} else if (name == "load_from") {
|
2014-02-10 02:10:30 +01:00
|
|
|
generate(p_value);
|
|
|
|
} else
|
|
|
|
return false;
|
|
|
|
|
|
|
|
return true;
|
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
bool PHashTranslation::_get(const StringName &p_name, Variant &r_ret) const {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
|
|
|
String name = p_name.operator String();
|
2017-03-05 16:44:50 +01:00
|
|
|
if (name == "hash_table")
|
|
|
|
r_ret = hash_table;
|
|
|
|
else if (name == "bucket_table")
|
|
|
|
r_ret = bucket_table;
|
|
|
|
else if (name == "strings")
|
|
|
|
r_ret = strings;
|
2014-02-10 02:10:30 +01:00
|
|
|
else
|
|
|
|
return false;
|
|
|
|
|
|
|
|
return true;
|
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
StringName PHashTranslation::get_message(const StringName &p_src_text) const {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
|
|
|
int htsize = hash_table.size();
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
if (htsize == 0)
|
2014-02-10 02:10:30 +01:00
|
|
|
return StringName();
|
|
|
|
|
|
|
|
CharString str = p_src_text.operator String().utf8();
|
2017-03-05 16:44:50 +01:00
|
|
|
uint32_t h = hash(0, str.get_data());
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2020-02-17 22:06:54 +01:00
|
|
|
const int *htr = hash_table.ptr();
|
2017-03-05 16:44:50 +01:00
|
|
|
const uint32_t *htptr = (const uint32_t *)&htr[0];
|
2020-02-17 22:06:54 +01:00
|
|
|
const int *btr = bucket_table.ptr();
|
2017-03-05 16:44:50 +01:00
|
|
|
const uint32_t *btptr = (const uint32_t *)&btr[0];
|
2020-02-17 22:06:54 +01:00
|
|
|
const uint8_t *sr = strings.ptr();
|
2017-03-05 16:44:50 +01:00
|
|
|
const char *sptr = (const char *)&sr[0];
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
uint32_t p = htptr[h % htsize];
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
if (p == 0xFFFFFFFF) {
|
2014-02-10 02:10:30 +01:00
|
|
|
return StringName(); //nothing
|
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
const Bucket &bucket = *(const Bucket *)&btptr[p];
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
h = hash(bucket.func, str.get_data());
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
int idx = -1;
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
for (int i = 0; i < bucket.size; i++) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
if (bucket.elem[i].key == h) {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
idx = i;
|
2014-02-10 02:10:30 +01:00
|
|
|
break;
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
if (idx == -1) {
|
2014-02-10 02:10:30 +01:00
|
|
|
return StringName();
|
|
|
|
}
|
|
|
|
|
|
|
|
if (bucket.elem[idx].comp_size == bucket.elem[idx].uncomp_size) {
|
|
|
|
|
|
|
|
String rstr;
|
2017-03-05 16:44:50 +01:00
|
|
|
rstr.parse_utf8(&sptr[bucket.elem[idx].str_offset], bucket.elem[idx].uncomp_size);
|
2014-02-10 02:10:30 +01:00
|
|
|
|
|
|
|
return rstr;
|
|
|
|
} else {
|
|
|
|
|
|
|
|
CharString uncomp;
|
2017-03-05 16:44:50 +01:00
|
|
|
uncomp.resize(bucket.elem[idx].uncomp_size + 1);
|
2017-11-25 04:07:54 +01:00
|
|
|
smaz_decompress(&sptr[bucket.elem[idx].str_offset], bucket.elem[idx].comp_size, uncomp.ptrw(), bucket.elem[idx].uncomp_size);
|
2014-02-10 02:10:30 +01:00
|
|
|
String rstr;
|
|
|
|
rstr.parse_utf8(uncomp.get_data());
|
|
|
|
return rstr;
|
|
|
|
}
|
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
void PHashTranslation::_get_property_list(List<PropertyInfo> *p_list) const {
|
2014-02-10 02:10:30 +01:00
|
|
|
|
Variant: Added 64-bit packed arrays, renamed Variant::REAL to FLOAT.
- Renames PackedIntArray to PackedInt32Array.
- Renames PackedFloatArray to PackedFloat32Array.
- Adds PackedInt64Array and PackedFloat64Array.
- Renames Variant::REAL to Variant::FLOAT for consistency.
Packed arrays are for storing large amount of data and creating stuff like
meshes, buffers. textures, etc. Forcing them to be 64 is a huge waste of
memory. That said, many users requested the ability to have 64 bits packed
arrays for their games, so this is just an optional added type.
For Variant, the float datatype is always 64 bits, and exposed as `float`.
We still have `real_t` which is the datatype that can change from 32 to 64
bits depending on a compile flag (not entirely working right now, but that's
the idea). It affects math related datatypes and code only.
Neither Variant nor PackedArray make use of real_t, which is only intended
for math precision, so the term is removed from there to keep only float.
2020-02-24 19:20:53 +01:00
|
|
|
p_list->push_back(PropertyInfo(Variant::PACKED_INT32_ARRAY, "hash_table"));
|
|
|
|
p_list->push_back(PropertyInfo(Variant::PACKED_INT32_ARRAY, "bucket_table"));
|
2020-02-17 22:06:54 +01:00
|
|
|
p_list->push_back(PropertyInfo(Variant::PACKED_BYTE_ARRAY, "strings"));
|
2017-03-05 16:44:50 +01:00
|
|
|
p_list->push_back(PropertyInfo(Variant::OBJECT, "load_from", PROPERTY_HINT_RESOURCE_TYPE, "Translation", PROPERTY_USAGE_EDITOR));
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
void PHashTranslation::_bind_methods() {
|
|
|
|
|
2017-08-09 13:19:41 +02:00
|
|
|
ClassDB::bind_method(D_METHOD("generate", "from"), &PHashTranslation::generate);
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|
|
|
|
|
2017-03-05 16:44:50 +01:00
|
|
|
PHashTranslation::PHashTranslation() {
|
2014-02-10 02:10:30 +01:00
|
|
|
}
|