### 迭代器
* 所有的容器都是可迭代的（iterable）

In [1]:
def is_iterable(param):
    try: 
        iter(param) 
        return True
    except TypeError:
        return False

params = [
    1234,
    '1234',
    [1, 2, 3, 4],
    set([1, 2, 3, 4]),
    {1:1, 2:2, 3:3, 4:4},
    (1, 2, 3, 4)
]
    
for param in params:
    print('{} is iterable? {}'.format(param, is_iterable(param)))

########## 输出 ##########


1234 is iterable? False
1234 is iterable? True
[1, 2, 3, 4] is iterable? True
{1, 2, 3, 4} is iterable? True
{1: 1, 2: 2, 3: 3, 4: 4} is iterable? True
(1, 2, 3, 4) is iterable? True


### 生成器
* 生成器是懒人版本的迭代器。
* 所有的容器都是可迭代的（iterable）。
* 生成器在 Python 的写法是用小括号括起来，(i for i in range(100000000))，即初始化了一个生成器。

In [3]:
import os
import psutil

# 显示当前 python 程序占用的内存大小
def show_memory_info(hint):
    pid = os.getpid()
    p = psutil.Process(pid)
    
    info = p.memory_full_info()
    memory = info.uss / 1024. / 1024
    print('{} memory used: {} MB'.format(hint, memory))

show_memory_info("show_info")    

show_info memory used: 53.125 MB


In [4]:
def test_iterator():
    show_memory_info('initing iterator')
    # 迭代器，消耗大量内存
    list_1 = [i for i in range(100000000)]
    show_memory_info('after iterator initiated')
    print(sum(list_1))
    show_memory_info('after sum called')

def test_generator():
    show_memory_info('initing generator')
    # 生成器，仅在被使用的时候才会被调用
    list_2 = (i for i in range(100000000))
    show_memory_info('after generator initiated')
    print(sum(list_2))
    show_memory_info('after sum called')

%time test_iterator()
%time test_generator()

########## 输出 ##########

initing iterator memory used: 53.1171875 MB
after iterator initiated memory used: 3888.265625 MB
4999999950000000
after sum called memory used: 3888.265625 MB
CPU times: total: 1.95 s
Wall time: 6.1 s
initing generator memory used: 54.26171875 MB
after generator initiated memory used: 54.26171875 MB
4999999950000000
after sum called memory used: 54.26171875 MB
CPU times: total: 1.31 s
Wall time: 4.47 s


In [8]:
def generator(k):
    i = 1
    while True:
        yield i ** k
        i += 1

gen_1 = generator(1)
gen_3 = generator(3)

print(gen_1)
print(gen_3)
print()

def get_sum(n):
    sum_1, sum_3 = 0, 0
    for i in range(n):
        # next_1 在yield处停止
        next_1 = next(gen_1)
        ## next_3 从 yield处，继续执行
        next_3 = next(gen_3)
        print('next_1 = {}, next_3 = {}'.format(next_1, next_3))
        sum_1 += next_1
        sum_3 += next_3
    print(sum_1 * sum_1, sum_3)

get_sum(8)

########## 输出 ##########

<generator object generator at 0x00000179CA83BC60>
<generator object generator at 0x00000179CA83BD30>

next_1 = 1, next_3 = 1
next_1 = 2, next_3 = 8
next_1 = 3, next_3 = 27
next_1 = 4, next_3 = 64
next_1 = 5, next_3 = 125
next_1 = 6, next_3 = 216
next_1 = 7, next_3 = 343
next_1 = 8, next_3 = 512
1296 1296


In [9]:
# 不使用迭代器的情况
def index_normal(L, target):
    result = []
    for i, num in enumerate(L):
        if num == target:
            result.append(i)
    return result

print(index_normal([1, 6, 2, 4, 5, 2, 8, 6, 3, 2], 2))

########## 输出 ##########


[2, 5, 9]


In [12]:
# 使用生成器来实现上面的功能
def index_generator(L, target):
    for i, num in enumerate(L):
        if num == target:
            yield i

print(list(index_generator([1, 6, 2, 4, 5, 2, 8, 6, 3, 2], 2)))

########## 输出 ##########

[2, 5, 9]


### 判断一个列表是否是另一个列表的子序列（元素不必连续出现）
* 使用贪心算法
* 使用python的生成器实现

In [15]:
def is_subsequence(a, b):
    b = iter(b)
    return all(i in b for i in a)

print(is_subsequence([1, 3, 5], [1, 2, 3, 4, 5]))
print(is_subsequence([1, 4, 3], [1, 2, 3, 4, 5]))

########## 输出 ##########


True
False


In [None]:
### 代码解释
def is_subsequence(a, b):
    b = iter(b)
    print(b)

    gen = (i for i in a)
    print(gen)

    for i in gen:
        print(i)

    gen = ((i in b) for i in a)
    print(gen)

    for i in gen:
        print(i)

    return all(((i in b) for i in a))

print(is_subsequence([1, 3, 5], [1, 2, 3, 4, 5]))
print(is_subsequence([1, 4, 3], [1, 2, 3, 4, 5]))

########## 输出 ##########


<list_iterator object at 0x00000179CA836BC0>
<generator object is_subsequence.<locals>.<genexpr> at 0x000001790AA175E0>
1
3
5
<generator object is_subsequence.<locals>.<genexpr> at 0x000001790AD3F300>
True
True
True
False
<list_iterator object at 0x000001790AD572E0>
<generator object is_subsequence.<locals>.<genexpr> at 0x000001790AA17C60>
1
4
3
<generator object is_subsequence.<locals>.<genexpr> at 0x000001790AD3F760>
True
True
False
False


In [16]:
b = (i for i in range(5))

print(2 in b)
print(4 in b)
print(3 in b)

########## 输出 ##########

True
True
False
