string ="""
<html>
<head>
<title>The Dormouse's story</title>
</head>
<body>
<p class="title" name="dromouse">
<b>The Dormouse's story</b>
</p>
<p class="story">
Once upon a time there were three little sisters; and their names were
<a class="mysis" href="http://example.com/elsie" id="link1">
<b>the first b tag<b>
Elsie
</a>,
<a class="mysis" href="http://example.com/lacie" id="link2" myname="kong">
Lacie
</a>and
<a class="mysis" href="http://example.com/tillie" id="link3">
Tillie
</a>;and they lived at the bottom of a well.
</p>
<p class="story">
myStory
<a>the end a tag</a>
</p>
<a>the p tag sibling</a>
</body>
</html>
"""
soup = BeautifulSoup(string,'lxml')
"""
=============================================================
选择器返回的永远是列表,元素是soup对象,对象有两种方法:
1、soup.attrs 返回属性
2、soup.text 返回文本
=============================================================
"""# 标签选择器
result = soup.select('a b')# 类选择器
result = soup.select('.mysis')# id选择器
result = soup.select('#link2')# 属性选择器
result = soup.select('a[id="link3"]')# 组合选择器
result = soup.select('.story #link2')# 获取属性
result = soup.select('.story #link2')[0].attrs # 注意这里attrs方法要针对bs元素,不能是列表# 获取文本
result = soup.select('.story #link2')[0].text # 注意这里text方法要针对bs元素,不能是列表